Mock Version: 6.1 Mock Version: 6.1 Mock Version: 6.1 ENTER ['do_with_status'](['bash', '--login', '-c', '/usr/bin/rpmbuild -bs --noclean --target x86_64 --nodeps /builddir/build/SPECS/llama-cpp.spec'], chrootPath='/var/lib/mock/f43-build-59176177-6571523/root'env={'TERM': 'vt100', 'SHELL': '/bin/bash', 'HOME': '/builddir', 'HOSTNAME': 'mock', 'PATH': '/usr/bin:/bin:/usr/sbin:/sbin', 'PROMPT_COMMAND': 'printf "\\033]0;\\007"', 'PS1': ' \\s-\\v\\$ ', 'LANG': 'C.UTF-8'}shell=Falselogger=timeout=201600uid=1000gid=425user='mockbuild'unshare_net=TrueprintOutput=Falsenspawn_args=['--capability=cap_ipc_lock', '--bind=/tmp/mock-resolv.yq32gdp_:/etc/resolv.conf', '--bind=/dev/btrfs-control', '--bind=/dev/mapper/control', '--bind=/dev/fuse', '--bind=/dev/loop-control', '--bind=/dev/loop0', '--bind=/dev/loop1', '--bind=/dev/loop2', '--bind=/dev/loop3', '--bind=/dev/loop4', '--bind=/dev/loop5', '--bind=/dev/loop6', '--bind=/dev/loop7', '--bind=/dev/loop8', '--bind=/dev/loop9', '--bind=/dev/loop10', '--bind=/dev/loop11']) Using nspawn with args ['--capability=cap_ipc_lock', '--bind=/tmp/mock-resolv.yq32gdp_:/etc/resolv.conf', '--bind=/dev/btrfs-control', '--bind=/dev/mapper/control', '--bind=/dev/fuse', '--bind=/dev/loop-control', '--bind=/dev/loop0', '--bind=/dev/loop1', '--bind=/dev/loop2', '--bind=/dev/loop3', '--bind=/dev/loop4', '--bind=/dev/loop5', '--bind=/dev/loop6', '--bind=/dev/loop7', '--bind=/dev/loop8', '--bind=/dev/loop9', '--bind=/dev/loop10', '--bind=/dev/loop11'] Executing command: ['/usr/bin/systemd-nspawn', '-q', '-M', '72348ef9be984a4ba6d8bd994fe97f0e', '-D', '/var/lib/mock/f43-build-59176177-6571523/root', '-a', '-u', 'mockbuild', '--capability=cap_ipc_lock', '--bind=/tmp/mock-resolv.yq32gdp_:/etc/resolv.conf', '--bind=/dev/btrfs-control', '--bind=/dev/mapper/control', '--bind=/dev/fuse', '--bind=/dev/loop-control', '--bind=/dev/loop0', '--bind=/dev/loop1', '--bind=/dev/loop2', '--bind=/dev/loop3', '--bind=/dev/loop4', '--bind=/dev/loop5', '--bind=/dev/loop6', '--bind=/dev/loop7', '--bind=/dev/loop8', '--bind=/dev/loop9', '--bind=/dev/loop10', '--bind=/dev/loop11', '--console=pipe', '--setenv=TERM=vt100', '--setenv=SHELL=/bin/bash', '--setenv=HOME=/builddir', '--setenv=HOSTNAME=mock', '--setenv=PATH=/usr/bin:/bin:/usr/sbin:/sbin', '--setenv=PROMPT_COMMAND=printf "\\033]0;\\007"', '--setenv=PS1= \\s-\\v\\$ ', '--setenv=LANG=C.UTF-8', '--resolv-conf=off', 'bash', '--login', '-c', '/usr/bin/rpmbuild -bs --noclean --target x86_64 --nodeps /builddir/build/SPECS/llama-cpp.spec'] with env {'TERM': 'vt100', 'SHELL': '/bin/bash', 'HOME': '/builddir', 'HOSTNAME': 'mock', 'PATH': '/usr/bin:/bin:/usr/sbin:/sbin', 'PROMPT_COMMAND': 'printf "\\033]0;\\007"', 'PS1': ' \\s-\\v\\$ ', 'LANG': 'C.UTF-8', 'SYSTEMD_NSPAWN_TMPFS_TMP': '0', 'SYSTEMD_SECCOMP': '0'} and shell False Building target platforms: x86_64 Building for target x86_64 setting SOURCE_DATE_EPOCH=1741478400 Wrote: /builddir/build/SRPMS/llama-cpp-b4580-2.fc43.src.rpm Child return code was: 0 ENTER ['do_with_status'](['bash', '--login', '-c', '/usr/bin/rpmbuild -bb --noclean --target x86_64 --nodeps /builddir/build/SPECS/llama-cpp.spec'], chrootPath='/var/lib/mock/f43-build-59176177-6571523/root'env={'TERM': 'vt100', 'SHELL': '/bin/bash', 'HOME': '/builddir', 'HOSTNAME': 'mock', 'PATH': '/usr/bin:/bin:/usr/sbin:/sbin', 'PROMPT_COMMAND': 'printf "\\033]0;\\007"', 'PS1': ' \\s-\\v\\$ ', 'LANG': 'C.UTF-8'}shell=Falselogger=timeout=201600uid=1000gid=425user='mockbuild'unshare_net=TrueprintOutput=Falsenspawn_args=['--capability=cap_ipc_lock', '--bind=/tmp/mock-resolv.yq32gdp_:/etc/resolv.conf', '--bind=/dev/btrfs-control', '--bind=/dev/mapper/control', '--bind=/dev/fuse', '--bind=/dev/loop-control', '--bind=/dev/loop0', '--bind=/dev/loop1', '--bind=/dev/loop2', '--bind=/dev/loop3', '--bind=/dev/loop4', '--bind=/dev/loop5', '--bind=/dev/loop6', '--bind=/dev/loop7', '--bind=/dev/loop8', '--bind=/dev/loop9', '--bind=/dev/loop10', '--bind=/dev/loop11']) Using nspawn with args ['--capability=cap_ipc_lock', '--bind=/tmp/mock-resolv.yq32gdp_:/etc/resolv.conf', '--bind=/dev/btrfs-control', '--bind=/dev/mapper/control', '--bind=/dev/fuse', '--bind=/dev/loop-control', '--bind=/dev/loop0', '--bind=/dev/loop1', '--bind=/dev/loop2', '--bind=/dev/loop3', '--bind=/dev/loop4', '--bind=/dev/loop5', '--bind=/dev/loop6', '--bind=/dev/loop7', '--bind=/dev/loop8', '--bind=/dev/loop9', '--bind=/dev/loop10', '--bind=/dev/loop11'] Executing command: ['/usr/bin/systemd-nspawn', '-q', '-M', 'bfb7186e0282442f80fe16f3765d5696', '-D', '/var/lib/mock/f43-build-59176177-6571523/root', '-a', '-u', 'mockbuild', '--capability=cap_ipc_lock', '--bind=/tmp/mock-resolv.yq32gdp_:/etc/resolv.conf', '--bind=/dev/btrfs-control', '--bind=/dev/mapper/control', '--bind=/dev/fuse', '--bind=/dev/loop-control', '--bind=/dev/loop0', '--bind=/dev/loop1', '--bind=/dev/loop2', '--bind=/dev/loop3', '--bind=/dev/loop4', '--bind=/dev/loop5', '--bind=/dev/loop6', '--bind=/dev/loop7', '--bind=/dev/loop8', '--bind=/dev/loop9', '--bind=/dev/loop10', '--bind=/dev/loop11', '--console=pipe', '--setenv=TERM=vt100', '--setenv=SHELL=/bin/bash', '--setenv=HOME=/builddir', '--setenv=HOSTNAME=mock', '--setenv=PATH=/usr/bin:/bin:/usr/sbin:/sbin', '--setenv=PROMPT_COMMAND=printf "\\033]0;\\007"', '--setenv=PS1= \\s-\\v\\$ ', '--setenv=LANG=C.UTF-8', '--resolv-conf=off', 'bash', '--login', '-c', '/usr/bin/rpmbuild -bb --noclean --target x86_64 --nodeps /builddir/build/SPECS/llama-cpp.spec'] with env {'TERM': 'vt100', 'SHELL': '/bin/bash', 'HOME': '/builddir', 'HOSTNAME': 'mock', 'PATH': '/usr/bin:/bin:/usr/sbin:/sbin', 'PROMPT_COMMAND': 'printf "\\033]0;\\007"', 'PS1': ' \\s-\\v\\$ ', 'LANG': 'C.UTF-8', 'SYSTEMD_NSPAWN_TMPFS_TMP': '0', 'SYSTEMD_SECCOMP': '0'} and shell False Building target platforms: x86_64 Building for target x86_64 setting SOURCE_DATE_EPOCH=1741478400 Executing(%mkbuilddir): /bin/sh -e /var/tmp/rpm-tmp.Jh7QWW Executing(%prep): /bin/sh -e /var/tmp/rpm-tmp.n4LXEy + umask 022 + cd /builddir/build/BUILD/llama-cpp-b4580-build + cd /builddir/build/BUILD/llama-cpp-b4580-build + rm -rf llama.cpp-b4580 + /usr/lib/rpm/rpmuncompress -x /builddir/build/SOURCES/llama.cpp-b4580.tar.gz + STATUS=0 + '[' 0 -ne 0 ']' + cd llama.cpp-b4580 + /usr/bin/chmod -Rf a+rX,u+w,g-w,o-w . + sed -i -e 's/POSITION_INDEPENDENT_CODE ON/POSITION_INDEPENDENT_CODE ON SOVERSION b4580/' src/CMakeLists.txt + sed -i -e 's/POSITION_INDEPENDENT_CODE ON/POSITION_INDEPENDENT_CODE ON SOVERSION b4580/' ggml/src/CMakeLists.txt + sed -i '/target_link_libraries(ggml-hip PRIVATE ggml-base.*/aset_target_properties(ggml-hip PROPERTIES SOVERSION b4580)' ggml/src/ggml-hip/CMakeLists.txt + sed -i '/target_compile_features(${GGML_CPU_NAME} PRIVATE c_std_11.*/aset_target_properties(${GGML_CPU_NAME} PROPERTIES SOVERSION b4580)' ggml/src/ggml-cpu/CMakeLists.txt + sed -i '/#include ' src/llama-mmap.h + rm -rf exmples/llma.android + find . -name .gitignore -exec rm -rf '{}' ';' + RPM_EC=0 ++ jobs -p + exit 0 Executing(%build): /bin/sh -e /var/tmp/rpm-tmp.Uh1y9e + umask 022 + cd /builddir/build/BUILD/llama-cpp-b4580-build + CFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer ' + export CFLAGS + CXXFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer' + export CXXFLAGS + FFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -I/usr/lib64/gfortran/modules ' + export FFLAGS + FCFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -I/usr/lib64/gfortran/modules ' + export FCFLAGS + VALAFLAGS=-g + export VALAFLAGS + RUSTFLAGS='-Copt-level=3 -Cdebuginfo=2 -Ccodegen-units=1 -Cstrip=none -Cforce-frame-pointers=yes -Clink-arg=-specs=/usr/lib/rpm/redhat/redhat-package-notes --cap-lints=warn' + export RUSTFLAGS + LDFLAGS='-Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes ' + export LDFLAGS + LT_SYS_LIBRARY_PATH=/usr/lib64: + export LT_SYS_LIBRARY_PATH + CC=hipcc + export CC + CXX=hipcc + export CXX + cd llama.cpp-b4580 + CFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer ' + export CFLAGS + CXXFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer' + export CXXFLAGS + FFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -I/usr/lib64/gfortran/modules ' + export FFLAGS + FCFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -I/usr/lib64/gfortran/modules ' + export FCFLAGS + VALAFLAGS=-g + export VALAFLAGS + RUSTFLAGS='-Copt-level=3 -Cdebuginfo=2 -Ccodegen-units=1 -Cstrip=none -Cforce-frame-pointers=yes -Clink-arg=-specs=/usr/lib/rpm/redhat/redhat-package-notes --cap-lints=warn' + export RUSTFLAGS + LDFLAGS='-Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes ' + export LDFLAGS + LT_SYS_LIBRARY_PATH=/usr/lib64: + export LT_SYS_LIBRARY_PATH + CC=hipcc + export CC + CXX=hipcc + export CXX + /usr/bin/cmake -S . -B redhat-linux-build -DCMAKE_C_FLAGS_RELEASE:STRING=-DNDEBUG -DCMAKE_CXX_FLAGS_RELEASE:STRING=-DNDEBUG -DCMAKE_Fortran_FLAGS_RELEASE:STRING=-DNDEBUG -DCMAKE_VERBOSE_MAKEFILE:BOOL=ON -DCMAKE_INSTALL_DO_STRIP:BOOL=OFF -DCMAKE_INSTALL_PREFIX:PATH=/usr -DCMAKE_INSTALL_FULL_SBINDIR:PATH=/usr/bin -DCMAKE_INSTALL_SBINDIR:PATH=bin -DINCLUDE_INSTALL_DIR:PATH=/usr/include -DLIB_INSTALL_DIR:PATH=/usr/lib64 -DSYSCONF_INSTALL_DIR:PATH=/etc -DSHARE_INSTALL_PREFIX:PATH=/usr/share -DLIB_SUFFIX=64 -DBUILD_SHARED_LIBS:BOOL=ON -DCMAKE_INSTALL_LIBDIR=lib64 -DCMAKE_SKIP_RPATH=ON -DGGML_AVX=OFF -DGGML_AVX2=OFF -DGGML_AVX512=OFF -DGGML_AVX512_VBMI=OFF -DGGML_AVX512_VNNI=OFF -DGGML_FMA=OFF -DGGML_F16C=OFF -DGGML_HIP=ON '-DAMDGPU_TARGETS=gfx900;gfx906:xnack-;gfx908:xnack-;gfx90a:xnack+;gfx90a:xnack-;gfx942;gfx1010;gfx1012;gfx1030;gfx1031;gfx1035;gfx1100;gfx1101;gfx1102;gfx1103;gfx1150;gfx1151;gfx1152;gfx1200;gfx1201' -DLLAMA_BUILD_EXAMPLES=OFF -DLLAMA_BUILD_TESTS=OFF -- The C compiler identification is Clang 19.0.0 -- The CXX compiler identification is Clang 19.0.0 -- Detecting C compiler ABI info -- Detecting C compiler ABI info - done -- Check for working C compiler: /usr/bin/hipcc - skipped -- Detecting C compile features -- Detecting C compile features - done -- Detecting CXX compiler ABI info -- Detecting CXX compiler ABI info - done -- Check for working CXX compiler: /usr/bin/hipcc - skipped -- Detecting CXX compile features -- Detecting CXX compile features - done -- Found Git: /usr/bin/git (found version "2.49.0") fatal: not a git repository (or any of the parent directories): .git fatal: not a git repository (or any of the parent directories): .git sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory -- Setting GGML_NATIVE_DEFAULT to OFF -- Performing Test CMAKE_HAVE_LIBC_PTHREAD -- Performing Test CMAKE_HAVE_LIBC_PTHREAD - Success -- Found Threads: TRUE -- Warning: ccache not found - consider installing it for faster compilation or disable this warning with GGML_CCACHE=OFF -- CMAKE_SYSTEM_PROCESSOR: x86_64 -- Including CPU backend -- Could NOT find OpenMP_C (missing: OpenMP_C_FLAGS OpenMP_C_LIB_NAMES) -- Could NOT find OpenMP_CXX (missing: OpenMP_CXX_FLAGS OpenMP_CXX_LIB_NAMES) -- Could NOT find OpenMP (missing: OpenMP_C_FOUND OpenMP_CXX_FOUND) CMake Warning at ggml/src/ggml-cpu/CMakeLists.txt:54 (message): OpenMP not found Call Stack (most recent call first): ggml/src/CMakeLists.txt:312 (ggml_add_cpu_backend_variant_impl) -- x86 detected -- Adding CPU backend variant ggml-cpu: -msse4.2 GGML_SSE42 CMake Warning at ggml/src/ggml-hip/CMakeLists.txt:27 (message): Setting hipcc as the C++ compiler is legacy behavior. Prefer setting the HIP compiler directly. See README for details. CMake Warning (dev) at /usr/lib64/cmake/hip/hip-config-amd.cmake:70 (message): AMDGPU_TARGETS is deprecated. Please use GPU_TARGETS instead. Call Stack (most recent call first): /usr/lib64/cmake/hip/hip-config.cmake:159 (include) ggml/src/ggml-hip/CMakeLists.txt:39 (find_package) This warning is for project developers. Use -Wno-dev to suppress it. -- Performing Test HIP_CLANG_SUPPORTS_PARALLEL_JOBS -- Performing Test HIP_CLANG_SUPPORTS_PARALLEL_JOBS - Success -- HIP and hipBLAS found -- Including HIP backend fatal: not a git repository (or any of the parent directories): .git fatal: not a git repository (or any of the parent directories): .git CMake Warning at common/CMakeLists.txt:32 (message): Git repository not found; to enable automatic generation of build info, make sure Git is installed and the project is a Git repository. -- Configuring done (19.0s) -- Generating done (0.1s) CMake Warning: Manually-specified variables were not used by the project: CMAKE_Fortran_FLAGS_RELEASE CMAKE_INSTALL_DO_STRIP INCLUDE_INSTALL_DIR LIB_INSTALL_DIR LIB_SUFFIX SHARE_INSTALL_PREFIX SYSCONF_INSTALL_DIR -- Build files have been written to: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build + /usr/bin/cmake --build redhat-linux-build -j6 --verbose Change Dir: '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' Run Build Command(s): /usr/bin/cmake -E env VERBOSE=1 /usr/bin/gmake -f Makefile -j6 /usr/bin/cmake -S/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 -B/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build --check-build-system CMakeFiles/Makefile.cmake 0 /usr/bin/cmake -E cmake_progress_start /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/CMakeFiles /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build//CMakeFiles/progress.marks /usr/bin/gmake -f CMakeFiles/Makefile2 all gmake[1]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f ggml/src/CMakeFiles/ggml-base.dir/build.make ggml/src/CMakeFiles/ggml-base.dir/depend /usr/bin/gmake -f common/CMakeFiles/build_info.dir/build.make common/CMakeFiles/build_info.dir/depend gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/CMakeFiles/ggml-base.dir/DependInfo.cmake "--color=" [ 0%] Generating build details from Git cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 && /usr/bin/cmake -DMSVC= -DCMAKE_C_COMPILER_VERSION=19.0.0 -DCMAKE_C_COMPILER_ID=Clang -DCMAKE_VS_PLATFORM_NAME= -DCMAKE_C_COMPILER=/usr/bin/hipcc -P /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/cmake/build-info-gen-cpp.cmake gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f ggml/src/CMakeFiles/ggml-base.dir/build.make ggml/src/CMakeFiles/ggml-base.dir/build gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 1%] Building C object ggml/src/CMakeFiles/ggml-base.dir/ggml.c.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu11 -fPIC -Wshadow -Wstrict-prototypes -Wpointer-arith -Wmissing-prototypes -Werror=implicit-int -Werror=implicit-function-declaration -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wdouble-promotion -MD -MT ggml/src/CMakeFiles/ggml-base.dir/ggml.c.o -MF CMakeFiles/ggml-base.dir/ggml.c.o.d -o CMakeFiles/ggml-base.dir/ggml.c.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml.c [ 1%] Building CXX object ggml/src/CMakeFiles/ggml-base.dir/ggml-opt.cpp.o [ 2%] Building CXX object ggml/src/CMakeFiles/ggml-base.dir/ggml-threading.cpp.o [ 3%] Building CXX object ggml/src/CMakeFiles/ggml-base.dir/ggml-backend.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT ggml/src/CMakeFiles/ggml-base.dir/ggml-opt.cpp.o -MF CMakeFiles/ggml-base.dir/ggml-opt.cpp.o.d -o CMakeFiles/ggml-base.dir/ggml-opt.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-opt.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT ggml/src/CMakeFiles/ggml-base.dir/ggml-backend.cpp.o -MF CMakeFiles/ggml-base.dir/ggml-backend.cpp.o.d -o CMakeFiles/ggml-base.dir/ggml-backend.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-backend.cpp -- Found Git: /usr/bin/git (found version "2.49.0") cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT ggml/src/CMakeFiles/ggml-base.dir/ggml-threading.cpp.o -MF CMakeFiles/ggml-base.dir/ggml-threading.cpp.o.d -o CMakeFiles/ggml-base.dir/ggml-threading.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-threading.cpp [ 4%] Building C object ggml/src/CMakeFiles/ggml-base.dir/ggml-alloc.c.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu11 -fPIC -Wshadow -Wstrict-prototypes -Wpointer-arith -Wmissing-prototypes -Werror=implicit-int -Werror=implicit-function-declaration -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wdouble-promotion -MD -MT ggml/src/CMakeFiles/ggml-base.dir/ggml-alloc.c.o -MF CMakeFiles/ggml-base.dir/ggml-alloc.c.o.d -o CMakeFiles/ggml-base.dir/ggml-alloc.c.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-alloc.c fatal: not a git repository (or any of the parent directories): .git fatal: not a git repository (or any of the parent directories): .git sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common/CMakeFiles/build_info.dir/DependInfo.cmake "--color=" gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f common/CMakeFiles/build_info.dir/build.make common/CMakeFiles/build_info.dir/build gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 5%] Building CXX object common/CMakeFiles/build_info.dir/build-info.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/build_info.dir/build-info.cpp.o -MF CMakeFiles/build_info.dir/build-info.cpp.o.d -o CMakeFiles/build_info.dir/build-info.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/build-info.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 6%] Building C object ggml/src/CMakeFiles/ggml-base.dir/ggml-quants.c.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu11 -fPIC -Wshadow -Wstrict-prototypes -Wpointer-arith -Wmissing-prototypes -Werror=implicit-int -Werror=implicit-function-declaration -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wdouble-promotion -MD -MT ggml/src/CMakeFiles/ggml-base.dir/ggml-quants.c.o -MF CMakeFiles/ggml-base.dir/ggml-quants.c.o.d -o CMakeFiles/ggml-base.dir/ggml-quants.c.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-quants.c sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 7%] Building CXX object ggml/src/CMakeFiles/ggml-base.dir/gguf.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_base_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT ggml/src/CMakeFiles/ggml-base.dir/gguf.cpp.o -MF CMakeFiles/ggml-base.dir/gguf.cpp.o.d -o CMakeFiles/ggml-base.dir/gguf.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/gguf.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 7%] Built target build_info [ 8%] Linking CXX shared library ../../bin/libggml-base.so cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/cmake -E cmake_link_script CMakeFiles/ggml-base.dir/link.txt --verbose=1 sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory clang++: warning: argument unused during compilation: '-Xarch_host -fstack-protector-strong' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-Xarch_host -fcf-protection' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-specs=/usr/lib/rpm/redhat/redhat-package-notes' [-Wunused-command-line-argument] /usr/bin/hipcc -fPIC -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -Xlinker --dependency-file=CMakeFiles/ggml-base.dir/link.d -Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes -shared -Wl,-soname,libggml-base.so.b4580 -o ../../bin/libggml-base.so.b4580 "CMakeFiles/ggml-base.dir/ggml.c.o" "CMakeFiles/ggml-base.dir/ggml-alloc.c.o" "CMakeFiles/ggml-base.dir/ggml-backend.cpp.o" "CMakeFiles/ggml-base.dir/ggml-opt.cpp.o" "CMakeFiles/ggml-base.dir/ggml-threading.cpp.o" "CMakeFiles/ggml-base.dir/ggml-quants.c.o" "CMakeFiles/ggml-base.dir/gguf.cpp.o" -lm cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/cmake -E cmake_symlink_library ../../bin/libggml-base.so.b4580 ../../bin/libggml-base.so.b4580 ../../bin/libggml-base.so gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 8%] Built target ggml-base /usr/bin/gmake -f ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/build.make ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/depend /usr/bin/gmake -f ggml/src/CMakeFiles/ggml-cpu.dir/build.make ggml/src/CMakeFiles/ggml-cpu.dir/depend gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/CMakeFiles/ggml-cpu.dir/DependInfo.cmake "--color=" gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/DependInfo.cmake "--color=" gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/build.make ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/build gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f ggml/src/CMakeFiles/ggml-cpu.dir/build.make ggml/src/CMakeFiles/ggml-cpu.dir/build gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 8%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.cpp.o [ 10%] Building C object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-quants.c.o [ 11%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/acc.cu.o [ 11%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-aarch64.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/ggml-cpu.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/acc.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/acc.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/acc.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu11 -fPIC -Wshadow -Wstrict-prototypes -Wpointer-arith -Wmissing-prototypes -Werror=implicit-int -Werror=implicit-function-declaration -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wdouble-promotion -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-quants.c.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-quants.c.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-quants.c.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/ggml-cpu-quants.c [ 12%] Building C object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.c.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-aarch64.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-aarch64.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-aarch64.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/ggml-cpu-aarch64.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu11 -fPIC -Wshadow -Wstrict-prototypes -Wpointer-arith -Wmissing-prototypes -Werror=implicit-int -Werror=implicit-function-declaration -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wdouble-promotion -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.c.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.c.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.c.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/ggml-cpu.c [ 13%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-hbm.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-hbm.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-hbm.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-hbm.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/ggml-cpu-hbm.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 14%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-traits.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-traits.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-traits.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-traits.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/ggml-cpu-traits.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ [ 14%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/amx.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/amx.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/amx.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/amx.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/amx/amx.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ [ 15%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/mmq.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/mmq.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/mmq.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/mmq.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/amx/mmq.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 16%] Building CXX object ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/llamafile/sgemm.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_SSE42 -DGGML_USE_CPU_AARCH64 -DGGML_USE_LLAMAFILE -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_cpu_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -msse4.2 -MD -MT ggml/src/CMakeFiles/ggml-cpu.dir/ggml-cpu/llamafile/sgemm.cpp.o -MF CMakeFiles/ggml-cpu.dir/ggml-cpu/llamafile/sgemm.cpp.o.d -o CMakeFiles/ggml-cpu.dir/ggml-cpu/llamafile/sgemm.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cpu/llamafile/sgemm.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory 6 warnings generated when compiling for gfx1012. [ 17%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/arange.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/arange.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/arange.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/arange.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu [ 17%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/argmax.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/argmax.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/argmax.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/argmax.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ [ 18%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/argsort.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/argsort.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/argsort.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/argsort.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ [ 19%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/binbcast.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/binbcast.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/binbcast.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/binbcast.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu [ 20%] Linking CXX shared library ../../bin/libggml-cpu.so cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/cmake -E cmake_link_script CMakeFiles/ggml-cpu.dir/link.txt --verbose=1 sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory clang++: warning: argument unused during compilation: '-Xarch_host -fstack-protector-strong' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-Xarch_host -fcf-protection' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-specs=/usr/lib/rpm/redhat/redhat-package-notes' [-Wunused-command-line-argument] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx1031. 7 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1100. /usr/bin/hipcc -fPIC -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -Xlinker --dependency-file=CMakeFiles/ggml-cpu.dir/link.d -Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes -shared -Wl,-soname,libggml-cpu.so.b4580 -o ../../bin/libggml-cpu.so.b4580 "CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.c.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu.cpp.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-aarch64.cpp.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-hbm.cpp.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-quants.c.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/ggml-cpu-traits.cpp.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/amx.cpp.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/amx/mmq.cpp.o" "CMakeFiles/ggml-cpu.dir/ggml-cpu/llamafile/sgemm.cpp.o" ../../bin/libggml-base.so.b4580 cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/cmake -E cmake_symlink_library ../../bin/libggml-cpu.so.b4580 ../../bin/libggml-cpu.so.b4580 ../../bin/libggml-cpu.so gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 20%] Built target ggml-cpu [ 21%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/clamp.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/clamp.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/clamp.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/clamp.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1101. 7 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | 6 struct { | ^ warnings generated when compiling for gfx1150. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. 7 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1200. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx1200. 7 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1200. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx900. 7 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx90a. 7 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. 7 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/acc.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for host. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ [ 22%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/concat.cu.o /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/concat.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/concat.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/concat.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct {6 | ^ warnings generated when compiling for gfx908. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1200. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 9 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/arange.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for host. [ 22%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/conv-transpose-1d.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/conv-transpose-1d.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/conv-transpose-1d.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/conv-transpose-1d.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu 6 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. 7 warnings generated when compiling for gfx1201. 9 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cu:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argmax.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for host. [ 23%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/convert.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/convert.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/convert.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/convert.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx942. 9 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx906. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 17 warnings generated when compiling for gfx1012. 7 warnings generated when compiling for gfx900. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/argsort.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ 298 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for host. [ 24%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/count-equal.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/count-equal.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/count-equal.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/count-equal.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx908. 9 warnings generated when compiling for gfx1031. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx906. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ 7 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 6 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src91_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 7 warnings generated when compiling for gfx908. 9 warnings generated when compiling for gfx1100. 17 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 6 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ 7 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx1100. 9 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/clamp.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 6 warnings generated when compiling for host. [ 25%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/cpy.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/cpy.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/cpy.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/cpy.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ 7 warnings generated when compiling for gfx1031. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ 17 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1031. 9 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 7 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ 17 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 9 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 7 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 17 warnings generated when compiling for gfx1103. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 9 warnings generated when compiling for gfx1150. 7 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 7 warnings generated when compiling for gfx1101. 13 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 17 warnings generated when compiling for gfx1150. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/binbcast.cu:359:11: warning: 'break' will never be executed [-Wunreachable-code-break] 359 | } break; | ^~~~~ 9 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for host. [ 26%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/cross-entropy-loss.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/cross-entropy-loss.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/cross-entropy-loss.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/cross-entropy-loss.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 7 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 17 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ 9 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 7 warnings generated when compiling for gfx1103. 17 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ 6 warnings generated when compiling for gfx1031. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ 9 warnings generated when compiling for gfx1200. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 6 warnings generated when compiling for gfx1012. 13 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const 7 warnings generated when compiling for gfx1150. int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 9 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ 6 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | 7 warnings generated when compiling for gfx1151. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h 281 | struct { | ^ :281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | s/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ truct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 9 warnings generated when compiling for gfx900. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, con7 warnings generated when compiling for gfx1152. st int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 6 warnings generated when compiling for gfx1031. 17 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h struct { | ^ :213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | str/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ uct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ 281 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1150. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ 7 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx1035. 17 warnings generated when compiling for gfx906. 9 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struc/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ t { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | st/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ ruct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 17 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1100. 7 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ 13 warnings generated when compiling for gfx1151. 9 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 17 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1101. 7 warnings generated when compiling for gfx900. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ : warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.ht:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ * x = (src_t *) vx; /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 9 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 6 warnings generated when compiling for gfx1102. 7 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 9 warnings generated when compiling for gfx90a. 17 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 6 warnings generated when compiling for gfx1103. 13 warnings generated when compiling for gfx1200. 7 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:33: warning: unused parameter 'p0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:4:47: warning: unused parameter 'd0' [-Wunused-parameter] 4 | const int s0, const int p0, const int d0, const int output_size, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:5:79: warning: unused parameter 'src0_ne3' [-Wunused-parameter] 5 | const int src0_ne0, const int src0_ne1, const int src0_ne2, const int src0_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:39: warning: unused parameter 'src1_ne1' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:59: warning: unused parameter 'src1_ne2' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:6:79: warning: unused parameter 'src1_ne3' [-Wunused-parameter] 6 | const int src1_ne0, const int src1_ne1, const int src1_ne2, const int src1_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:38: warning: unused parameter 'dst_ne1' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:57: warning: unused parameter 'dst_ne2' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:7:76: warning: unused parameter 'dst_ne3' [-Wunused-parameter] 7 | const int dst_ne0, const int dst_ne1, const int dst_ne2, const int dst_ne3, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:78:19: warning: unused variable 'kernel_size' [-Wunused-variable] 78 | const int64_t kernel_size = ggml_nelements(src0); | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/conv-transpose-1d.cu:79:19: warning: unused variable 'input_size' [-Wunused-variable] 79 | const int64_t input_size = ggml_nelements(src1); | ^~~~~~~~~~ 17 warnings generated when compiling for host. [ 27%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/diagmask.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/diagmask.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/diagmask.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/diagmask.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ 7 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ 6 warnings generated when compiling for gfx1150. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 9 warnings generated when compiling for gfx942. 13 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:41:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 41 | if (blockIdx.y < ne01) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:67:20: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 67 | if (blockIdx.z < ne02) { // src0 | ~~~~~~~~~~ ^ ~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/concat.cu:218:17: warning: 'break' will never be executed [-Wunreachable-code-break] 218 | break; | ^~~~~ 9 warnings generated when compiling for host. [ 27%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ 13 warnings generated when compiling for gfx900. 7 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ 254 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h 213 | struct { | ^ :281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/count-equal.cu:62:13: warning: 'break' will never be executed [-Wunreachable-code-break] 62 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 7 warnings generated when compiling for host. [ 28%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f32.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f32.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f32.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f32.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu 6 warnings generated when compiling for gfx1200. 13 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ :213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { 13 warnings generated when compiling for gfx908. | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 30 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx906. 13 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ 30 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | In file included from struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 6 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 6 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 13 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. 30 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *__restrict' to 'type-parameter-0-0 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:467:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 467 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:33:31: warning: cast from 'const void *' to 'int *' drops const qualifier [-Wcast-qual] 33 | const int * x0 = ((int *) vx) + blockIdx.x * nint; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:470:9: note: in instantiation of function template specialization 'dequantize_block_q8_0_f16' requested here 470 | dequantize_block_q8_0_f16<<>>(vx, y, k); | ^ 6 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'float *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:635:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 635 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to '__half *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary<__half, float>' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:682:20: note: in instantiation of function template specialization 'convert_unary_cuda<__half, float>' requested here 682 | return convert_unary_cuda; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:580:33: warning: cast from 'const void *' to 'hip_bfloat16 *' drops const qualifier [-Wcast-qual] 580 | const src_t * x = (src_t *) vx; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:588:5: note: in instantiation of function template specialization 'convert_unary' requested here 588 | convert_unary<<>>(vx, y, k); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/convert.cu:684:20: note: in instantiation of function template specialization 'convert_unary_cuda' requested here 684 | return convert_unary_cuda; | ^ 13 warnings generated when compiling for host. [ 29%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ 6 warnings generated when compiling for gfx1151. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cross-entropy-loss.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for host. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ [ 30%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/getrows.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/getrows.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/getrows.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/getrows.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 80 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx942. 80 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/cpy.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. 30 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for host. [ 31%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/ggml-cuda.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/ggml-cuda.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/ggml-cuda.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/ggml-cuda.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1030. 80 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. 7 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 29 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel,In file included from nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 80 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx906. 7 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 29 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 7 warnings generated when compiling for gfx1035. 80 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. 29 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 30 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ 7 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 30 warnings generated when compiling for gfx1031. 80 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 29 warnings generated when compiling for gfx1031. 7 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 80 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 7 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/diagmask.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 29 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 6 warnings generated when compiling for host. [ 31%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/gla.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/gla.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/gla.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/gla.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ 7 warnings generated when compiling for gfx1103. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 29 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 30 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx1103. 7 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 29 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 7 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 6 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 29 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 7 warnings generated when compiling for gfx1152. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 29 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 30 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 29 warnings generated when compiling for gfx1150. 7 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 7 warnings generated when compiling for gfx900. 29 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. 80 warnings generated when compiling for gfx1201. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 7 warnings generated when compiling for gfx906. 29 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ const int ne2, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ 7 warnings generated when compiling for gfx908. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 30 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 29 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ 7 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 80 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ 6 warnings generated when compiling for gfx1101. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 29 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ 7 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 29 warnings generated when compiling for gfx900. 30 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 80 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 30 warnings generated when compiling for gfx1102. 29 warnings generated when compiling for gfx906. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/getrows.cu:201:13: warning: 'break' will never be executed [-Wunreachable-code-break] 201 | break; | ^~~~~ 7 warnings generated when compiling for host. [ 32%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/im2col.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/im2col.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/im2col.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/im2col.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ 80 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1010. 29 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 80 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ 6 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 30 warnings generated when compiling for gfx1103. 29 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:6: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:7: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:145:13: warning: 'break' will never be executed [-Wunreachable-code-break] 145 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:118:17: warning: 'break' will never be executed [-Wunreachable-code-break] 118 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:90:17: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:67:21: warning: 'break' will never be executed [-Wunreachable-code-break] 67 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:42:21: warning: 'break' will never be executed [-Wunreachable-code-break] 42 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn.cu:141:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_wmma_f16_case<256, 32, __half>' requested here 141 | ggml_cuda_flash_attn_ext_wmma_f16_case<256, cols_per_block, half>(ctx, dst); | ^ 6 warnings generated when compiling for gfx1030. 80 warnings generated when compiling for host. [ 33%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmq.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmq.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmq.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmq.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ 6 warnings generated when compiling for gfx1151. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 29 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1152. 29 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:5: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:22: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: warning: variable length arrays in C++ are a Clang extension [-Wvla-cxx-extension] 132 | char archName[archLen + 1]; | ^~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:132:19: note: read of non-const variable 'archLen' is not allowed in a constant expression /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:131:9: note: declared here 131 | int archLen = strlen(devName); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1150. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:52: warning: unused parameter 'buffer' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2837:67: warning: unused parameter 'size' [-Wunused-parameter] 2837 | bool ggml_backend_cuda_register_host_buffer(void * buffer, size_t size) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3142:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3142 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3137:13: warning: 'break' will never be executed [-Wunreachable-code-break] 3137 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3134:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3134 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3125:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3125 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3118:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3118 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3113:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3113 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3108:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3108 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3103:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3103 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3061:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3061 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3057:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3057 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:3040:15: warning: 'break' will never be executed [-Wunreachable-code-break] 3040 | } break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/ggml-cuda.cu:2978:13: warning: 'break' will never be executed [-Wunreachable-code-break] 2978 | break; | ^~~~~ 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 29 warnings generated when compiling for host. [ 34%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmv.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmv.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmv.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmv.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ 15 warnings generated when compiling for gfx1031. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ 6 warnings generated when compiling for gfx1103. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1012. 30 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 15 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 6 warnings generated when compiling for gfx1151. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx1152. 15 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx1200. 15 warnings generated when compiling for gfx1150. 30 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ 15 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 6 warnings generated when compiling for gfx900. 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx90a. 30 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | 6 warnings generated when compiling for gfx906. const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 15 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ 6 warnings generated when compiling for gfx908. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/gla.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for host. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct [ 35%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmvq.cu.o { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hcd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmvq.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmvq.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmvq.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu :281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | 6 warnings generated when compiling for gfx1102. const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx90a. 15 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ nrows_x, nrows_y, nrows_ds/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ t); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | m/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ tmp[ncols_y][row/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hs:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ _per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ 182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ mul_mat_vec_q<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h : 298 : 9 : warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] mul_mat _298v | e c _ q < t y p es,t r5u>c>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_ma/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.ht_:v192e:c9_:q << <192b | l o c k _ n u m ss,t rbulcotc k{_ d i| m ^s , 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restric/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cut_:_80 :x48,: cwarning: osuggest braces around initialization of subobject [-Wmissing-braces]n st int * __restrict _80_ | y , ffllooaatt *t m_p_[rnecsotlrsi_cyt]_[_r oswusm_,p ecro_ncsutd ai_nbtl o&c kk]0 0=) {{0 . 0| f ^} ; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f};/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] _y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 1704 | const int * __restrict__ x, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ :2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu6 warnings generated when compiling for gfx90a. :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[15 warnings generated when compiling for gfx906. ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 30 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx942. 15 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 30 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/im2col.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for host. [ 36%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/norm.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/norm.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/norm.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/norm.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu 15 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 6 warnings generated when compiling for gfx1010. 15 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ 281 | struct { | ^/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ :298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ struct { | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h ^ :298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx1012. 30 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. 30 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cu:90:13: warning: 'break' will never be executed [-Wunreachable-code-break] 90 | break; | ^~~~~ 15 warnings generated when compiling for host. [ 36%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/opt-step-adamw.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/opt-step-adamw.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/opt-step-adamw.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/opt-step-adamw.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx906. 30 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmv.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for host. [ 37%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/out-prod.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/out-prod.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/out-prod.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/out-prod.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 293 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrIn file included from ows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ 80 | float tmp[ncols_y][rows_per_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuhcuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ on_no_fattn_vec_case/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu(const int D) { | ^ :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ const int ne03, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | co/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ nst int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int n/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cue13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ :126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (roIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ ws_per_cuda_block == 1 |/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ | row0 + threadIdx.x < nro/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hws_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ :213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < ro/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ ws_per_cuda_block && (rows_p/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ er_cuda_block == 1/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h :298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, 6 warnings generated when compiling for gfx900. ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ Idx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ *((uint32_t *) &KQ_max_scale)/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_f/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ attn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_m/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ ask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 704 | flash_attn_combi/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ ne_results | ^/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ :13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h :192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ mul_mat_vec_q<<>>(vx, vy/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | s/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cutruct { | ^ :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cu6 warnings generated when compiling for gfx1030. da_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 30 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ launch_fattn(ctx, dst, fattn_kernel, nw/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ arps, cols_per_block, true,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ :281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/norm.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ t *) &KQ_max_scale) &= /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:319:13: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<16, 4, false>' requested here 319 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:293:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 293 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:299:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 299 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f32.cu:344:9: note: in instantiation of function template specialization 'launch_fattn_tile_f32_64_128<32, 1, false>' requested here 344 | launch_fattn_tile_f32_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for host. [ 38%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/pad.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/pad.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/pad.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/pad.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu 6 warnings generated when compiling for host. [ 39%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/pool2d.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/pool2d.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/pool2d.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/pool2d.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cu:2: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/opt-step-adamw.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 6 warnings generated when compiling for host. 8 warnings generated when compiling for gfx1012. [ 40%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/quantize.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/quantize.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/quantize.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/quantize.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu 6 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ :192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ :254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &K/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuhQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ :1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 8 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 15 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 6 warnings generated when compiling for gfx1151. 15 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 8 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:24:19: warning: unused parameter 'ne00' [-Wunused-parameter] 24 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:27:19: warning: unused parameter 'ne03' [-Wunused-parameter] 27 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:28:19: warning: unused parameter 'ne10' [-Wunused-parameter] 28 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:31:19: warning: unused parameter 'ne13' [-Wunused-parameter] 31 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:32:19: warning: unused parameter 'ne31' [-Wunused-parameter] 32 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:33:19: warning: unused parameter 'nb31' [-Wunused-parameter] 33 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:36:19: warning: unused parameter 'nb03' [-Wunused-parameter] 36 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:39:19: warning: unused parameter 'nb13' [-Wunused-parameter] 39 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:40:19: warning: unused parameter 'nb21' [-Wunused-parameter] 40 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:41:19: warning: unused parameter 'nb22' [-Wunused-parameter] 41 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:42:19: warning: unused parameter 'nb23' [-Wunused-parameter] 42 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:43:19: warning: unused parameter 'ne0' [-Wunused-parameter] 43 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:44:19: warning: unused parameter 'ne1' [-Wunused-parameter] 44 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:45:19: warning: unused parameter 'ne2' [-Wunused-parameter] 45 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:46:19: warning: unused parameter 'ne3' [-Wunused-parameter] 46 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:323:13: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<16, 4, false>' requested here 323 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:294:13: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 294 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:300:13: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 300 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/fattn-tile-f16.cu:348:9: note: in instantiation of function template specialization 'launch_fattn_tile_f16_64_128<32, 1, false>' requested here 348 | launch_fattn_tile_f16_64_128(ctx, dst); | ^ 8 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 30 warnings generated when compiling for host. [ 41%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/rope.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/rope.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/rope.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/rope.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 8 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 15 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 15 warnings generated when compiling for gfx1101. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ 6 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1152. 15 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 8 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx1200. 15 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/out-prod.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for host. [ 41%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/scale.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/scale.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/scale.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/scale.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 6 warnings generated when compiling for gfx1030. 15 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h9:: 281warning: :anonymous types declared in an anonymous union are an extension [-Wnested-anon-types]9 : warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | 281 | s tstructr {u c t| ^{ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 8 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 8 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 6 warnings generated when compiling for gfx1031. 15 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 6 warnings generated when compiling for gfx1030. 8 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 6 warnings generated when compiling for gfx1031. 8 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 6 warnings generated when compiling for gfx1035. 8 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 15 warnings generated when compiling for gfx906. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 6 warnings generated when compiling for gfx1100. 8 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pool2d.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 15 warnings generated when compiling for gfx908. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:56: warning: comparison of integers of different signs: 'unsigned int' and 'int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/pad.cu:17:35: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 17 | if (nidx < ne00 && blockIdx.y < ne01 && blockIdx.z < ne02*ne03) { | ~~~~~~~~~~ ^ ~~~~ 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1100. 8 warnings generated when compiling for host. [ 42%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/softmax.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/softmax.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/softmax.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/softmax.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu 6 warnings generated when compiling for host. [ 43%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/sum.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/sum.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/sum.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/sum.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1010. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 6 warnings generated when compiling for gfx1103. 15 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 293 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 6 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ 15 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cuh:4: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/quantize.cu:167:13: warning: 'break' will never be executed [-Wunreachable-code-break] 167 | break; | ^~~~~ 15 warnings generated when compiling for host. [ 44%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/sumrows.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/sumrows.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/sumrows.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/sumrows.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 161 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y,In file included from nr/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuo:w1s: _In file included from d/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuhs:t1,: In file included from s/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuht:r20e: am/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h);: 171 :| 9 ^: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: 191 | warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] mul_ m213a | t _ v e c _ q < tsytpreu,c t6 >{< < <| b ^l ock_nums, block_dims, 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<< > > ( v xs,t rvuyc,t d{s t ,| ^n cols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 6 warnings generated when compiling for gfx1102. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/scale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for host. [ 45%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/tsembd.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/tsembd.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/tsembd.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/tsembd.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. 161 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_6 warnings generated when compiling for gfx942. vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h>:(171v:x9,: vwarning: y,anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] dst, ncols_x, 171n | r o w s _ x , nsrtorwusc_ty ,{ n r| o ^w s_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu { : 80| ^:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sum.cu:10: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for host. [ 45%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/unary.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/unary.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/unary.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/unary.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. 6 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx906. 6 warnings generated when compiling for gfx1103. 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/sumrows.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx908. 6 warnings generated when compiling for gfx1151. 6 warnings generated when compiling for host. [ 46%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/upscale.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/upscale.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/upscale.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/upscale.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ 6 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1152. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 7 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 161 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for gfx1200. 6 warnings generated when compiling for gfx90a. 7 warnings generated when compiling for gfx1012. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ 6 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst,In file included from ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][ro/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hws_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ ncols_y][rows_per/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ _cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, v213y | , d s t , n csotlrsu_cxt, {n r o| w ^s _x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1101. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for gfx1201. 7 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for gfx1102. 7 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/softmax.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for host. 6 warnings generated when compiling for gfx90a. [ 47%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/wkv6.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/wkv6.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/wkv6.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/wkv6.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for gfx1103. 7 warnings generated when compiling for gfx1035. 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 7 warnings generated when compiling for gfx1100. 6 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for gfx1012. 7 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/rope.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 7 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx1030. 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1152. 6 warnings generated when compiling for host. [ 48%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 7 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1031. 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 161 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/tsembd.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1010. 7 warnings generated when compiling for gfx1150. 6 warnings generated when compiling for host. [ 49%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu 6 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx1201. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx900. 64 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 61 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 7 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_k/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cue:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ rnel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1030. 7 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 61 warnings generated when compiling for gfx1012. 6 warnings generated when compiling for gfx1102. 6 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 7 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct In file included from {/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1103. 61 warnings generated when compiling for gfx1030. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 7 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ 213 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h| ^ :192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h 281 | struct { | ^ :213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5:In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_att/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ n_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 7 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1035. 61 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for gfx942. 6 warnings generated when compiling for gfx1151. 7 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/unary.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1100. 161 warnings generated when compiling for gfx1101. 61 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 6 warnings generated when compiling for host. [ 50%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu 7 warnings generated when compiling for gfx90a. 6 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh :554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ dst, ncols_x, nrows_x, nrows_y, nrows_dst)/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ ; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ (vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ >>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ uda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy,In file included from d/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cus:t,3 : nIn file included from c/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuho:l2s: _/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhx:,554 :n24r:o wwarning: scast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual]_ x, nrows_y, nr o554w | s _ d s t ) ; *| ( ^( uint3/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu2:_237t: 5*:) note: &in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested hereK Q_max_scale) & =237 | f t z _ mmauslk_;m a t| _ ^v ec_q_cuda' requested here_ Q8_0>(vx ,704 | v y , dfslta,s hn_caotltsn__xc,o mnbrionwes__rxe,s unlrtosw | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ _dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>In file included from >/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ (vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 197 | mul_mat_vec_q<<>>(In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ rows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_In file included from p/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ er_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<7 warnings generated when compiling for gfx90a. >>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: In file included from note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu :3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24 :300 | warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] mul_mat_vec_q_cu d554a | < G G M L _ T Y P*E(_(IuQi3n_tX3X2S_>t( v*x), &vKyQ,_ mdasxt_,s cnacloel)s _&x=, fntrzo_wmsa_sxk,; n r| o ^w s_y, ncols_y,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh :n704r:o5w:s _note: din instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested heres t, stream )704; | | ^ flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_yIn file included from , n/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cuc:o3l: sIn file included from _/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuhy:,2 : n/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhro:w554s:_24d:s twarning: ,cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] stream); | ^ 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncIn file included from ols/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu_:y3]: [In file included from r/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuho:w2s: _/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhp:e554r:_24c:u dwarning: acast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual]_ block] = {0.0f}; 554 | | ^~~~ | { } *((uint32_t /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu*:)179 :&13K:Q _note: min instantiation of function template specialization 'mul_mat_vec_q' requested herea x_scale) &= ftz_mas k179; | | ^ mul/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh_:m704a:t5_:v enote: cin instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here_ q704< | < < b l ofclka_sh_antutmns_,c bolombicnke__dreismsu, 0l, tssl>e>l_(bvlxo,ck svy> , | ^d st, n/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuhc:o491:l9: snote: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here_ x, nrows _x491, nr | ow s _y , n ro w lsa_undcsh_fta)ttn<;D , p| a ^ rall/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:el_314b:5: note: lin instantiation of function template specialization 'mul_mat_vec_q_cuda' requested hereoc ks>(ctx, d s314t | , f at tn _kmeurnle_mlat,_v encw_aqr_pcsu,da(vxl, ovcyk,, tdrstu,e ncols_x,, ntrrouwes)_;x , | n ^r ows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh :1882 | : /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh : 554 : 24 : warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] mul_mat_vec_q < < < b*l(o(cuki_nntu3m2s_,t b*l)o c&kK_Qd_immasx,_ s0c,a lset)r e&a=m >f>t>z(_vmxa,s kv;y , | d ^s t, ncols_x, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhn:r704o:w5s:_ xnote: ,in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here nrows_y, n704r | o w s _ dfslta)s;h _ a| t ^t n_com/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cub:i314n:e5_:r enote: sin instantiation of function template specialization 'mul_mat_vec_q_cuda' requested hereu ltsu l _| m ^a t_vec/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh_:q505_:c5u:d anote: ' requested hereG GML_TYPE_IQ1_M> (505v | x , v yl,a udnscth,_ fnactotlns<_Dx,, pnarroawlsl_exl,_ bnlroocwkss_>y(,c tnxc,o ldss_ty,, fnartotwns__kdesrtn,e ls,t rnewaamr)p;s , | c ^ ols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, ds/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cut,: 80f:a48t:t nwarning: _suggest braces around initialization of subobject [-Wmissing-braces]k ernel, nwarps, cols_ p80e | r _ b l ofclko,a tt rtumep,[ ntcroules)_;y ] [| r ^o ws_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cudaIn file included from warning: (cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual]v x, vy, dst, ncol s554_ | x , n r o w s _*x(,( unirnotw3s2__yt, *n)c o&lKsQ__ym,a xn_rsocwasl_ed)s t&,= sfttrze_amma)s;k ; | ^| ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 64 warnings generated when compiling for gfx1101. 6 warnings generated when compiling for gfx1200. 61 warnings generated when compiling for gfx1100. 64 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 7 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 61 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/upscale.cu:22:37: warning: cast from 'const float *' to 'char *' drops const qualifier [-Wcast-qual] 22 | dst[index] = *(float *)((char *)x + i03 * nb03 + i02 * nb02 + i01 * nb01 + i00 * nb00); | ^ 64 warnings generated when compiling for gfx1012. 7 warnings generated when compiling for host. [ 50%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 64 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx906. 64 warnings generated when compiling for gfx1030. 64 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx908. 64 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 61 warnings generated when compiling for gfx1103. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1151. 61 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 6 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1030. 161 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 61 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 6 warnings generated when compiling for gfx942. 64 warnings generated when compiling for gfx1100. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_sca/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cul:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ e) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_m/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ ax_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | launch_fattn(ctx, dst, fattn_kerne/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cul:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ , nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ :554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ >(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn:(80c:t48x:, warning: dsuggest braces around initialization of subobject [-Wmissing-braces]s t, fattn_kernel, nw a80r | p s , cfollosa_tp etrm_pb[lnoccokl,s _tyr]u[er,o wtsr_upee)r;_ c u| d ^a _block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ _max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ , nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuIn file included from :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu :80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ el_blocks> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ : warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_ma/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cus:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ k; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_resu/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ lts | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | la/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ unch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ _attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 64 warnings generated when compiling for gfx1031. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/wkv6.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 6 warnings generated when compiling for host. [ 51%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1200. 61 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_tIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_att/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ n_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ 64 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 58 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1201. 61 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 58 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx1201. 64 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx1101. 58 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx900. 64 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ 491 | launch_fattn(ctx, dst, fattn_kernel, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ nwarps, cols_per_block, true, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fatt/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhn:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ (ctx, dst, fattn_kernel/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ , nwarps, cols_per_block, true, true); | ^ In file included from In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx1031. 64 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx906. 161 warnings generated when compiling for gfx1103. 64 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ const int nb01, | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter]/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ 38 | cons/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ t int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu*): 80&:48K:Q _warning: maxsuggest braces around initialization of subobject [-Wmissing-braces]_s cale) &= ftz_ ma80 | s k ; | ^ float t/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhm:p704[:n5c:o lnote: sin instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here_y][ rows_p er704_c | u d a _ bfllaoshc_attkn]_co mb=i n{e0_.re0sfu}l;t s <| D, ^~~~ p | { }a rallel_bl/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:o182c:13k: snote: >in instantiation of function template specialization 'mul_mat_vec_q' requested here | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: 182note: | in instantiation of function template specialization 'launch_fattn<96, 2>' requested here mul_ m491a | t _ v e c _ q < tlyapuen,c h3_>fs(,c t0x,, sdtsrte,a mf>a>t>t(nv_xk,e rvnye,l ,d sntw,a rnpcso,l sc_oxl,s _npreorw_sb_lxo,c kn,r otwrsu_ey,, tnrruoew)s;_ d s| t ^) ; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rowsIn file included from _/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cupe:r_3c: uIn file included from da/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh_:bl2oc: k/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh]: 554=: {240:.0f} ;warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] | ^~~~ | { } 554 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu : 191*:(13(:u inote: nin instantiation of function template specialization 'mul_mat_vec_q' requested heret 32_t *) &KQ_max_sc a191l | e ) & = f t z _ m a smku;l _ m| a ^t _vec_q5<:< ' requested herel ock_nums, 704b | l o c k _fdliamssh,_ a0t,t ns_tcroemabmi>n>e>_(rvexs,u lvtys,< Dd,s tp,a rnaclollesl__xb,l oncrkosw>s _ x| , ^ nrows_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuhy:,505 :n5r:o wnote: sin instantiation of function template specialization 'launch_fattn<96, 1>' requested here_ dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu :505272 | : 5 : note: lin instantiation of function template specialization 'mul_mat_vec_q_cuda' requested herea unch_fattnu(dcate(rvnxe,l ,v yn,w adrspts,, nccoollss__pxe,r _nbrloowcsk_,x ,t rnureo,w st_ryu,e )n;c o l| s ^_ y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vIn file included from x,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu :v3y: ,In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:ds2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuht:,554: nc24ol:s warning: _cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual]x , nrows_ x554 | , n *r(ow(su_iyn,t 3n2c_to *l)s &_KyQ_,m anxr_oswcsa_ldes)t &= f,t zs_tmraesam)k;; | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 64 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx908. 64 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_resIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ ults | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx90a. 64 warnings generated when compiling for gfx1150. 64 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, 58 warnings generated when compiling for gfx1101. | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx90a. 64 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx942. 64 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 58 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 61 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1152. 64 warnings generated when compiling for host. [ 52%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq1_s.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq1_s.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq1_s.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq1_s.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 64 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 58 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hIn file included from :213/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu::9:3 : warning: In file included from anonymous types declared in an anonymous union are an extension [-Wnested-anon-types]/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh :2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: 213warning: | cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] str u554c | t { | ^ *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ :1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const In file included from i/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ nt & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_sIn file included from t/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ ream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { 61 warnings generated when compiling for host. | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ [ 53%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_s.cu.o /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_s.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_s.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_s.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 161 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 58 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ 2805 | mul_mat_q_str/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ eam_k_fixup<<>> /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fIn file included from ixu/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cup:<3t: y/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuhpe:,14 :m35m:q _warning: xunused parameter 'Q' [-Wunused-parameter], MMQ_NWARPS, nee d14_ | c h e c k > < < :> >warning: unused parameter 'K' [-Wunused-parameter] | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh15: | 2891 : 13 : note: in instantiation of function template specialization 'launch_mul_mat_q' requested here const ch a2891r | * _ _ r e s t r i c tl_a_u nKc,h _ m| u ^l _mat_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuhq:<16t:y35p:e ,warning: unused parameter 'V' [-Wunused-parameter]1 12>(ctx ,16 | a r g s , s t recaomn)s;t c| h ^a r * __rest/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhr:i2691c:t16_:_ warning: Vcomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare], | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17 :352691: | warning: unused parameter 'mask' [-Wunused-parameter] if (17i | t ! = b l o cckoIndsxt. xc h|a|r j*t _!_=r ebsltorcikcItd_x_. ym)a s{k , | ~~ ^ ~~~~~~~~~~| ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hIn file included from :/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu281::39: :In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuhwarning: :2anonymous types declared in an anonymous union are an extension [-Wnested-anon-types]: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual]281 | struct 554{ | | ^ *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 64 warnings generated when compiling for gfx1201. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, nco/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ ls_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_blIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ ock] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ ims, 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ , 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13:In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested hereIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx,In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | 64 warnings generated when compiling for gfx908. ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 58 warnings generated when compiling for gfx1151. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ 64 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx90a. 58 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx906. 64 warnings generated when compiling for gfx90a. 58 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh: :In file included from 476/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh::92:: note: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuhin instantiation of function template specialization 'launch_fattn<96, 4>' requested here: 554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 476 | l554a | u n c h _ f a t t*n(<(Du,i npta3r2a_ltl e*l)_ b&lKoQc_kmsa>x(_cstcxa,l ed)s t&,= ffattzt_nm_aksekr;n e l| , ^ nwarps, cols/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh_:p704e:r5_:b lnote: oin instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested herec k, true, t704r | u e) ; f| l ^a sh_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 64 warnings generated when compiling for gfx908. 58 warnings generated when compiling for gfx1201. 64 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_coIn file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ mbine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, tIn file included from r/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ ue, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ = ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx900. 64 warnings generated when compiling for gfx90a. 64 warnings generated when compiling for host. [ 54%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 &In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 161 warnings generated when compiling for gfx1151. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 281358 warnings generated when compiling for gfx906. | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 64 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu::1263:: 98In file included from :/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh :warning: 1comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare]: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 126 | 171 | i f ( t hstrruecatd I{dx. x < | r ^ow s_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_qwarning: >>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadId/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hx.:x298 :<9 :n rwarning: oanonymous types declared in an anonymous union are an extension [-Wnested-anon-types]w s_dst)) { 298 | | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx908. 64 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<80, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<80, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<80, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<80, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<112, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<112, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<112, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<112, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx90a. 64 warnings generated when compiling for host. [ 54%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 58 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:14:35: warning: unused parameter 'Q' [-Wunused-parameter] 14 | const char * __restrict__ Q, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:15:35: warning: unused parameter 'K' [-Wunused-parameter] 15 | const char * __restrict__ K, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:16:35: warning: unused parameter 'V' [-Wunused-parameter] 16 | const char * __restrict__ V, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:17:35: warning: unused parameter 'mask' [-Wunused-parameter] 17 | const char * __restrict__ mask, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:18:35: warning: unused parameter 'dst' [-Wunused-parameter] 18 | float * __restrict__ dst, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:19:35: warning: unused parameter 'dst_meta' [-Wunused-parameter] 19 | float2 * __restrict__ dst_meta, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:20:21: warning: unused parameter 'scale' [-Wunused-parameter] 20 | const float scale, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:21:21: warning: unused parameter 'max_bias' [-Wunused-parameter] 21 | const float max_bias, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:22:21: warning: unused parameter 'm0' [-Wunused-parameter] 22 | const float m0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:23:21: warning: unused parameter 'm1' [-Wunused-parameter] 23 | const float m1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:24:24: warning: unused parameter 'n_head_log2' [-Wunused-parameter] 24 | const uint32_t n_head_log2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:25:21: warning: unused parameter 'logit_softcap' [-Wunused-parameter] 25 | const float logit_softcap, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:26:19: warning: unused parameter 'ne00' [-Wunused-parameter] 26 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:27:19: warning: unused parameter 'ne01' [-Wunused-parameter] 27 | const int ne01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:28:19: warning: unused parameter 'ne02' [-Wunused-parameter] 28 | const int ne02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:29:19: warning: unused parameter 'ne03' [-Wunused-parameter] 29 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:30:19: warning: unused parameter 'ne10' [-Wunused-parameter] 30 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:31:19: warning: unused parameter 'ne11' [-Wunused-parameter] 31 | const int ne11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:32:19: warning: unused parameter 'ne12' [-Wunused-parameter] 32 | const int ne12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:33:19: warning: unused parameter 'ne13' [-Wunused-parameter] 33 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:34:19: warning: unused parameter 'ne31' [-Wunused-parameter] 34 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:35:19: warning: unused parameter 'nb31' [-Wunused-parameter] 35 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:36:19: warning: unused parameter 'nb01' [-Wunused-parameter] 36 | const int nb01, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:37:19: warning: unused parameter 'nb02' [-Wunused-parameter] 37 | const int nb02, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:38:19: warning: unused parameter 'nb03' [-Wunused-parameter] 38 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:39:19: warning: unused parameter 'nb11' [-Wunused-parameter] 39 | const int nb11, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:40:19: warning: unused parameter 'nb12' [-Wunused-parameter] 40 | const int nb12, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:41:19: warning: unused parameter 'nb13' [-Wunused-parameter] 41 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:42:19: warning: unused parameter 'nb21' [-Wunused-parameter] 42 | const int nb21, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:43:19: warning: unused parameter 'nb22' [-Wunused-parameter] 43 | const int nb22, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:44:19: warning: unused parameter 'nb23' [-Wunused-parameter] 44 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:45:19: warning: unused parameter 'ne0' [-Wunused-parameter] 45 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:46:19: warning: unused parameter 'ne1' [-Wunused-parameter] 46 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:47:19: warning: unused parameter 'ne2' [-Wunused-parameter] 47 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:48:19: warning: unused parameter 'ne3' [-Wunused-parameter] 48 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<64, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<96, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<96, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<96, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<96, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<128, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:476:9: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 476 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 2>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:491:9: note: in instantiation of function template specialization 'launch_fattn<256, 2>' requested here 491 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-wmma-f16.cuh:505:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 505 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, true, true); | ^ 58 warnings generated when compiling for host. [ 55%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_s.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_s.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_s.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_s.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 78 warnings generated when compiling for gfx1031. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 293 warnings generated when compiling for gfx1201. 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[nco/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ ls_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ vx, vy, dst, ncols_x, nrows_x, nrows_y, nc/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ ols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ GGML_TYPE_Q4_0>(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ :2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ 2691 | if (it != blockIdx.x || j/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ t != blockIdx.y) { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh[:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, v/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ y, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhy:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ , nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(v/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhx:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ , vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh):2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ims, 0, stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nro/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691w:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ s_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_ds/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuht:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ <<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, nco/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ls_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh ^~~~:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ><<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh :2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhc:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ uda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cu/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ da_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 191 | mul_mat_vec_q<<' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ stream>>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. 78 warnings generated when compiling for gfx1035. 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h13:: note: 298in instantiation of function template specialization 'launch_mul_mat_q' requested here: 9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 2888 | 298 | l asutnrcuhc_tm u{l _ m| a ^t _q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 293 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. 293 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for gfx942. 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:57:34: warning: unused parameter 'nrows_x' [-Wunused-parameter] 57 | const int ncols_x, const int nrows_x, const int nrows_y, const int nrows_dst) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:418:13: warning: 'break' will never be executed [-Wunreachable-code-break] 418 | break; | ^~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:209:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 209 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:216:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 216 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:223:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 223 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:230:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 230 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:237:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 237 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:244:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 244 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:251:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 251 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:258:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 258 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<_>c>u d a| _ ^b lock] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh=: 2858{0:13.:0 fnote: }in instantiation of function template specialization 'launch_mul_mat_q' requested here; | ^~~~ | { } 2858 | l/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cua:u182n:c13h:_ mnote: uin instantiation of function template specialization 'mul_mat_vec_q' requested herel _mat_q(c t182x | , a r g s , s t r e ammu)l;_ m a| t ^_ vec_q16<:< => >b(lvoxc,k Ivdyx,. yd)s t{, n| c ~~ ^ ~~~~~~~~~~o ls_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(c/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cutx:, 80a:r48g:s ,warning: suggest braces around initialization of subobject [-Wmissing-braces]s tream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh80: | 2691 :16 : warning: fcomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] loat tmp[nco l2691s | _ y ] [ r o w s _ipfe r(_ictu d!a=_ bblloocckk]I d=x .x ||{ 0j.t0 f!}=; b l| o ^~~~c k I| d { }x .y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:265:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 265 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:272:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 272 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:279:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 279 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<> > | ^ if (/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuht:h2867r:e13a:d Inote: din instantiation of function template specialization 'launch_mul_mat_q' requested herex .x < rows _2867p | e r _ c u d a _ b l o c kl a&u&n c(hr_omwusl__pmeart__cqu=( c1t x|,| arrogws0, +s ttrheraema)d;I d x| . ^x < nrows_d/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhs:t2691):)16 :{ warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]| ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:286:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 286 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_m/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhat:_v2691e:c36_:q <<b>l>o(cvkxI,d xv.yx, |d|s tj,t n!c=o lbsl_oxc,k Indrxo.wys)_ x{, n| r ~~ ^ ~~~~~~~~~~o ws_y, nrows_dst/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh):;2805 : 9| : ^ note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 2805 | 293m | u l _ m amtu_lq__msattr_evaemc__kq__fciuxduap<Q(_vNxW,A RvPyS,, dnsete,d _ncchoelcsk_>x<,< m>)>; | | ^ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][r/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhow:s_2691p:e36r:_ cwarning: ucomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]d a_block] = {0.0f}; | ^~~~2691 | | { } if (it !/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu=: 191b:l13o:c knote: Iin instantiation of function template specialization 'mul_mat_vec_q' requested hered x.x || jt != blockI d191x | . y ) { | ~~ ^ ~~~~~~~~~~ mul_mat_v/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhe:c2813_:q9<:t ynote: pin instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested heree , 6><<>t>(vreaxm,_ kv_fiyx, udps ^< <' requested here_ xy_tilin g293, | b l o cmuk_l_mdaitm_sv,e c0,_ qstre_acm>uda><>G G M| L ^_ TYP/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhE:_2873:I13Q:2 _note: Sin instantiation of function template specialization 'launch_mul_mat_q' requested here> (vx, v 2873 | y, d s t , n co l s _x, lnraowsu_nx,c nrhow_sm_yu, lnc_olmas_ty_,q ea(cmt)x;, | ^ args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:293:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 293 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:300:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 300 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:307:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 307 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup80< | < < b l ofclko_antu mtsm_px[yn_ctoillsi_nyg],[ rbolwosc_kp_edri_mcsu,d a0_,b lsotcrke]a m=> >{>0 . 0| f ^} ; | ^~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh :| 2882 { }: 13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu :2882176 | : 13 : note: in instantiation of function template specialization 'mul_mat_vec_q' requested here launch_mul_ma t176_ | q < t y p e , 8 8 > (mcutlx_,m aatr_gvse,c _sqt| < ^< > > ( v x , ivfy ,( idts t!,= nbcloolcsk_Ixd,x .nxr o|w|s _jxt, !n=r obwlso_cyk,I dnxr.oyw)s _{d s t| ) ~~ ^ ~~~~~~~~~~; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu.x :|80:|48: jwarning: tsuggest braces around initialization of subobject [-Wmissing-braces] != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh: 281380: | 9 : note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested heref loat tmp[ncols_y][ r2813o | w s _ p e r _ c umdual__bmlaotc_kq]_ s=t r{e0a.m0_fk}_;f i x| u ^~~~p < t| y { }p e, mmq_x, MMQ_N/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuW:A179R:P13S:, note: nin instantiation of function template specialization 'mul_mat_vec_q' requested heree ed_check><<>><>< < b| l ^o ck_nu/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhm:s2882,: 13b:l onote: cin instantiation of function template specialization 'launch_mul_mat_q' requested herek _dims, 0, 2882s | t r e a m > > > ( v x , lvayu,n cdhs_tm,u ln_cmoalts__qx<,t ynpreo,w s _8x8,> (ncrtoxw,s _ayr,g sn,r oswtsr_edasmt));; | | ^ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:: 2691note: :in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here16 : warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 314 | 2691 | m u l _ m aitf_ v(eict_ q!_=c ubdlao (bvlxo,c kvIyd,x .dys)t ,{ n c| o ~~ ^ ~~~~~~~~~~l s_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:314:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 314 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_m/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cuul_:m80a:t48_:q (ctx, args ,80 | s t r e afml)o;a t | t ^m p[ncols_y]/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh[:r2691o:w16s:_ pwarning: ecomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]r _cuda_block ]2691 | = { 0 . 0 f } ;i f | ( ^~~~i t | ! { }= blockIdx.x ||/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu :j194t: 13!:= note: bin instantiation of function template specialization 'mul_mat_vec_q' requested herel ockIdx.y) { | ~~ ^ ~~~~~~~~~~ 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:321:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 321 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:328:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 328 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:176:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 176 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:179:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 179 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:182:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 182 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:185:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 185 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu(i:t126 :!98=: bwarning: lcomparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare]o ckIdx.x || jt != blockIdx.y )126 | { | ~~ ^ ~~~~~~~~~~ if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:188:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 188 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockI/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cudx.:y126): 98{: warning: | ~~ ^ ~~~~~~~~~~comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:191:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 191 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:194:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 194 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:80:48: warning: suggest braces around initialization of subobject [-Wmissing-braces] 80 | float tmp[ncols_y][rows_per_cuda_block] = {0.0f}; | ^~~~ | { } /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:197:13: note: in instantiation of function template specialization 'mul_mat_vec_q' requested here 197 | mul_mat_vec_q<<>>(vx, vy, dst, ncols_x, nrows_x, nrows_y, nrows_dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:335:5: note: in instantiation of function template specialization 'mul_mat_vec_q_cuda' requested here 335 | mul_mat_vec_q_cuda(vx, vy, dst, ncols_x, nrows_x, nrows_y, ncols_y, nrows_dst, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/mmvq.cu:126:98: warning: comparison of integers of different signs: 'unsigned int' and 'const int' [-Wsign-compare] 126 | if (threadIdx.x < rows_per_cuda_block && (rows_per_cuda_block == 1 || row0 + threadIdx.x < nrows_dst)) { | ~~~~~~~~~~~~~~~~~~ ^ ~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 293 warnings generated when compiling for host. [ 56%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ 78 warnings generated when compiling for gfx1152. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ 78 warnings generated when compiling for gfx908. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh: 2691:| 36 ^: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh :l194a:u105n:c hwarning: _function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn]m ul_mat_q194( | c t x , _a_rdgesv,i cset_r_e a_m_)f;o r c| e ^i nline__ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhv:o2691i:d16 :m mwarning: acomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]_ K8(const mm a2691_ | i n t _ A _ I 1 6iKf8 (&i tm m!a=_ Ab,l occoknIsdtx .mxm a|_|i njtt_ B!_=J 8bKl8o c&k Imdmxa._yB)) {{ | | ~~ ^ ~~~~~~~~~~ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq1_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 57%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for host. [ 58%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01In file included from ,/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ c/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ onst i/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hn:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ t & ne10, cons/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ t int & /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ ne1/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 1, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 59%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q2_k.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q2_k.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q2_k.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q2_k.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] 1022 | for (int k01 = 0; k01 < WARP_SIZE; k01 += QR2_K*VDR_Q2_K_Q8_1_MMQ) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 59%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q3_k.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q3_k.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q3_k.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q3_k.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 98 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_s.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 60%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] In file included from 281 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu : 3 : In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh :s3t: rIn file included from u/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuhct: 20{: | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h ^: 171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhlau:nc2691h_:m36u:l _mwarning: acomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]t_ q(ctx, args, stre am2691) | ; | ^ if /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh(:i2691t:16 :! =warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] blockIdx .2691x | | | j ti f !(i=t b!=l oblcokckIIddxx..x y|)| j{t ! = | bl ~~ ^ ~~~~~~~~~~o ckIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | i/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ f (it !=/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ :192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ if (it != blockIdx.x || jt != blockId/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ x.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ :298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<< >b>l o c| k ^I dx.y) /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh{: 2870 :| 13 ~~ ^ ~~~~~~~~~~: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh :28702813 | : 9 : note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | l a u nch _m u l_ mmualt__mqa_fi(xcutpx<,t yapreg,s ,m msqt_rxe,a mM);M Q | _ ^N WAR/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhP:S2691,: 16n:e ewarning: dcomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]_ chec k2691> | < < < b l o c ki_fn u(mits _!x=y _tbilloicngk,I dbx.lxo c|k|_ dijtm s!,= 0,b lostcrkeIamd>x.>y>) | { ^ | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh ~~ ^ ~~~~~~~~~~: 2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ype, mmq_x, MMQ_NWARPS, need_check><<' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ lock_dims, 0, stream>>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockI/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ dx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | i/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ f (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ :2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, con/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ st int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ :2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, str78eam); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 61%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_1.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_1.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_1.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_1.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<< > > i f| ^( it != /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhb:l2864o:c13k:I dnote: xin instantiation of function template specialization 'launch_mul_mat_q' requested here. x || jt ! =2864 | b l o c k I d x . y ) {l a u| n ~~ ^ ~~~~~~~~~~c h_mul_mat_qin instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here( ctx, args, stream) ;2813 | | ^ mul_ma/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuht:_2691q:_16s:t rwarning: ecomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]a m_k_fixupo>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh : 2691 : 36 : mwarning: ucomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]l _mat_q_stream_k_fixu p2691< | t y p e , m m qi_fx ,( iMtM Q!_=N WbAlRoPcSk,I dnxe.exd _|c|h ejctk >!<=< ' requested heres tream>>> | ^ 2813 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh : 2882 : 13 : note: min instantiation of function template specialization 'launch_mul_mat_q' requested hereu l_mat_q_st r2882e | a m _ k _ f i x u p < t ylpaeu,n cmhm_qm_ux,l _MmMaQt__NqW_(cchtexc,k >a>> | 2691 ^ | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh : 2867i:f13 :( inote: tin instantiation of function template specialization 'launch_mul_mat_q' requested here != blockI d2867x | . x | | j t ! = bllaoucnkcIhd_xm.uyl)_ m{a t _| q ~~ ^ ~~~~~~~~~~< type, 48>(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36 :2805 | warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] mul_mat_q_strea m2691_ | k _ f i x u p < tiyfp e(,i tm m!q=_ xb,l oMcMkQI_dNxW.AxR P|S|, jnte e!d=_ cbhleocckk>I' requested here, 0, stream>>> | ^ 2813 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh : 2879 : 13 :m unote: lin instantiation of function template specialization 'launch_mul_mat_q' requested here_ mat_q_stre a2879m | _ k _ f i x u p < t y p el,a umnmcqh__xm,u lM_MmQa_tN_WqAc(hcetcxk,> >> 2691| | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh :i2891f: 13(:i tnote: in instantiation of function template specialization 'launch_mul_mat_q' requested here! = blockIdx. x2891 | | | j t ! = b l o clkaIudnxc.hy_)m u{l _| ma ~~ ^ ~~~~~~~~~~t _q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(3c: tIn file included from x/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh,: 3a: rIn file included from g/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuhs:,20 : st/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hr:e171a:m9):; warning: | ^anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: 171comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] | str u2691c | t { | ^ if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x |/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh|:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ heck><<>> /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh| :2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != bloc/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhk:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ Idx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ d_check><<>> /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 2813 | mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ eed_check><<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ e, mmq_x, MMQ_NWARPS, need_check><<' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ims, 0, stream>>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh::2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 62%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_k.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_k.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_k.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_k.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 63%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockI/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhd:1053x:.y99) {: warning: | ~~ ^ ~~~~~~~~~~unused parameter 'k00' [-Wunused-parameter] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805: 91053: | note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here const int * __re s2805t | r i c t _ _ x ,m uclo_nmsatt _iqn_ts t*r e_a_mr_eks_tfriixcutp_<_t yyp,e ,f lmomaqt_ x*, _M_MrQe_sNtWrAiRcPtS_,_ nseuemd,_ cchoencskt> >> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q3_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 63%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_1.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_1.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_1.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_1.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 78 warnings generated when compiling for gfx942. 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 64%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_k.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_k.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_k.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_k.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ 78 warnings generated when compiling for gfx1101. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 2691 | if (it != blockIdx.x || jt != blockIdx.y) { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, st/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != block/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ Idx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ mmq_x, MMQ_NWARPS, need_check><<' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ _dims, 0, stream>>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhq:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ _stream_k_fixup<<' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ms_xy_tiling, block_dims, 0, stream>>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockI/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ dx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ e, mmq_x, MMQ_NWARPS, need_check><<' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ ck_dims, 0, stream>>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 65%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q6_k.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q6_k.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q6_k.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q6_k.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | cons/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuht in:t2691 *: 36:_ _rwarning: ecomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]s trict__ x, c2691on | s t i n t * _i_refs (tirt !i=ct b_l_o cykI,d xf.lxo ||at j t* _!_= brleostcrkiIdcxt__. syu)m { , | ~~ ^ ~~~~~~~~~~ const int/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh :&2805 :9k00:) note: { in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here| ^ 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. 78 warnings generated when compiling for gfx1152. 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_1.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 66%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q8_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q8_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q8_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q8_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 67%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. 60 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1012. 60 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhca:se2691_:i36m:p lwarning: y()c t{x , | d ~~ ^ ~~~~~~~~~~ st); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu: 32691: | In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh : 2 : /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh : 149i:f33 (:i twarning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare]! = blockI dx149. | x | | fjotr !(=i nbtl ok_cKkQI_d0x .y=) 0{; k| _ ~~ ^ ~~~~~~~~~~K Q_0 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh<: 2813D:/9s:i znote: ein instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested hereo f(in t2813) | ; k _ K Q _ 0m u+l_=m aWtA_RqP__sStIrZeEa)m _{k _ f| i ~~~~~~ ^ ~~~~~~~~~~~~~x up' requested here_ x, M M481Q | _ N W A R P S ,t ynpeeed__cKh e=c=k >GD>>> : | ^| ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh13::303 :note: 35:in instantiation of function template specialization 'launch_mul_mat_q' requested here note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 2873 | 303 | f a t tlna_uknecrhn_emlu_lt_ mfaatt_tqn<_tkyepren,e l 6=4> (fcltaxs,h _aargtst,n _svterce_aemx);t _ f1| 6 ^< D, /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhc:o2691l:s16_:p ewarning: rcomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]_ blo c2691k | , p a r a lilfe l(_bilto c!k=s ,b ltoypcek_IK,d xt.yxp e|_|V ,j tu s!e=_ lbolgoictk_Isdoxf.tyc)a p{> ; | ~~ ^ ~~~~~~~~~~| ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 60 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1100. 78 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh) {: 116| : ~~ ^ ~~~~~~~~~~37 : warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1102. 78 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] 1022 | for (int k01 = 0; k01 < WARP_SIZE; k01 += QR2_K*VDR_Q2_K_Q8_1_MMQ) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1022:5: warning: loop not unrolled: the optimizer was unable to perform the requested transformation; the transformation might be disabled or specified as part of an unsupported transformation ordering [-Wpass-failed=transform-warning] 60 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1035. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 98 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for host. [ 68%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. 78 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 60 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != block/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.hIdx:.y192): 9{: warning: | ~~ ^ ~~~~~~~~~~anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h :254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types]2879 | 254 | l a u n c h _smutlr_ucmta t{_ q <| t ^y pe, 80>(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: 2691warning: | anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] if (254i | t ! = b l o c ksItdrxu.cxt |{| j| t ^ != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 78 warnings generated when compiling for gfx90a. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1100. 78 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx900. 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q5_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 68%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1010. 60 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 60 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1035. 60 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx1201. 60 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx908. 60 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1150. 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx1201. 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps,In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ cols_per_block, need_f/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ 16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h_:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ flash_attn_ext_vec_f16_case_imp/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ l(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q4_0, GGML_TYPE_Q4_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for host. [ 69%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh_at:t2691n:_36e:x twarning: _vcomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]e c_f32_case_impl(ctx, dst); | /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh ^: 2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 2805 | m129u | l _ m a t _ q _ s t r e afmo_rk _(fiinxtu pi<0t y=p e0,; mim0q __ > > f| o ^r (i nt i/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh0: 2858=: 0;13: note: iin instantiation of function template specialization 'launch_mul_mat_q' requested here 0 < D /2858s | iz e of (i n t ) ; li0a +u=n WcAhR_mPu_SIl_ZEm)at_ {q < t| y ~~ ^ ~~~~~~~~~~~~~ pe, 24>(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1012. 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:333:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 333 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:343:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 343 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:346:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 346 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:356:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 356 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:359:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:369:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 369 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:372:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 372 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:384:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 384 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for host. [ 70%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1010. 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q6_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 71%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1012. 60 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 34 warnings generated when compiling for gfx1035. 78 warnings generated when compiling for host. [ 72%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx90a. 60 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<:<554<:b24l:o cwarning: kcast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual]_ nums_xy_tiling, block_d i554m | s, 0 , s t r e*a(m(>u>i>n t 3| 2 ^_ t *) &/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhK:Q2852_:m13a:x _note: sin instantiation of function template specialization 'launch_mul_mat_q' requested herec ale) &= f t2852z | _ m a s k ; | ^ launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_att/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuhn_v:ec2691_:e36x:t _warning: fcomparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare]3 2o;c k I| ^d x.y) { /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh| : ~~ ^ ~~~~~~~~~~311 :13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 311 | gg m2805l_ | c u d a _ f l a smhu_la_tmtant__eqx_ts_tvreeca_mf_3k2__fciaxsuep_c,( csttxr,e adms>t>)>; | | ^ ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuhnote: :in instantiation of function template specialization 'launch_mul_mat_q' requested here142 :33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 2897 | 142 | l a u n c h _ m u lf_omra t(_iqni(0c t' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1010. 34 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1031. 60 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1035. 34 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx1150. 34 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1151. 34 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 60 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 34 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1200. 34 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1103. 34 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 34 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1151. 60 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx900. 34 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impluint(32_ctt *x), &dKsQt)_;max _ s| c ^ale) &= ftz_mIn file included from as/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cuk;: 3 : In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh| : ^2 : /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx1201. 34 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 34 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx1200. 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ 34 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 128>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 128>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 128>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 128>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 128>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for host. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ [ 72%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu 60 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx900. 34 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scalIn file included from e) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 60 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 34 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 64>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 64>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 64>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 64>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 64>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx906. 34 warnings generated when compiling for host. [ 73%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 60 warnings generated when compiling for gfx90a. 34 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1010. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 34 warnings generated when compiling for gfx90a. 60 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:311:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 311 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:321:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 321 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:324:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 2, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 324 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:334:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 334 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:337:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 4, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 337 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:347:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 347 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:350:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 4, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 350 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:116:37: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 116 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:362:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_Q8_0, GGML_TYPE_Q8_0, true>' requested here 362 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:129:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 129 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:142:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 142 | for (int i0 = 0; i0 < D/sizeof(int); i0 += WARP_SIZE) { | ~~ ^ ~~~~~~~~~~~~~ 60 warnings generated when compiling for host. [ 74%] Building CXX object ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/hipcc -DGGML_BACKEND_BUILD -DGGML_BACKEND_SHARED -DGGML_HIP_NO_VMM -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_HIP -DUSE_PROF_API=1 -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -D__HIP_PLATFORM_AMD__=1 -Dggml_hip_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/.. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -x hip --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 -MD -MT ggml/src/ggml-hip/CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu.o -MF CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu.o.d -o CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu 34 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1010. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1012. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ 34 warnings generated when compiling for gfx942. /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx1030. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:483:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0<__half, 256>' requested here 483 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:482:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1<__half, 256>' requested here 482 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:481:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0<__half, 256>' requested here 481 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:480:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1<__half, 256>' requested here 480 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:479:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0<__half, 256>' requested here 479 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared<__half2>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:303:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f16<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 303 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f16; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:330:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 330 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:306:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 306 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f16.cuh:381:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f16_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 381 | ggml_cuda_flash_attn_ext_vec_f16_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for host. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 34 warnings generated when compiling for gfx1031. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1035. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1100. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 34 warnings generated when compiling for gfx1100. 34 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1101. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1103. 34 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1102. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 78 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 34 warnings generated when compiling for gfx1150. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ 34 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1103. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1151. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1150. 34 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 78 warnings generated when compiling for gfx90a. 34 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1151. 34 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1152. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx1200. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<64, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<64, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<64, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for host. 34 warnings generated when compiling for gfx1201. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx908. 34 warnings generated when compiling for gfx900. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx906. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx908. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx942. 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<128, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<128, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<128, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ 34 warnings generated when compiling for host. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx90a. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:1: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:563:47: warning: function 'on_no_fattn_vec_case' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 563 | static void on_no_fattn_vec_case(const int D) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:21:19: warning: unused parameter 'ne00' [-Wunused-parameter] 21 | const int ne00, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:24:19: warning: unused parameter 'ne03' [-Wunused-parameter] 24 | const int ne03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:25:19: warning: unused parameter 'ne10' [-Wunused-parameter] 25 | const int ne10, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:28:19: warning: unused parameter 'ne13' [-Wunused-parameter] 28 | const int ne13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:29:19: warning: unused parameter 'ne31' [-Wunused-parameter] 29 | const int ne31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:30:19: warning: unused parameter 'nb31' [-Wunused-parameter] 30 | const int nb31, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:33:19: warning: unused parameter 'nb03' [-Wunused-parameter] 33 | const int nb03, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:36:19: warning: unused parameter 'nb13' [-Wunused-parameter] 36 | const int nb13, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:39:19: warning: unused parameter 'nb23' [-Wunused-parameter] 39 | const int nb23, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:40:19: warning: unused parameter 'ne0' [-Wunused-parameter] 40 | const int ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:41:19: warning: unused parameter 'ne1' [-Wunused-parameter] 41 | const int ne1, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:42:19: warning: unused parameter 'ne2' [-Wunused-parameter] 42 | const int ne2, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:43:19: warning: unused parameter 'ne3' [-Wunused-parameter] 43 | const int ne3) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:247:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 247 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:494:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q8_0' requested here 494 | type_K == GGML_TYPE_Q8_0 ? vec_dot_fattn_vec_KQ_q8_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:196:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 196 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:493:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_1' requested here 493 | type_K == GGML_TYPE_Q5_1 ? vec_dot_fattn_vec_KQ_q5_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:149:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 149 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:492:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q5_0' requested here 492 | type_K == GGML_TYPE_Q5_0 ? vec_dot_fattn_vec_KQ_q5_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:105:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 105 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:491:36: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_1' requested here 491 | type_K == GGML_TYPE_Q4_1 ? vec_dot_fattn_vec_KQ_q4_1 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:65:33: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 65 | for (int k_KQ_0 = 0; k_KQ_0 < D/sizeof(int); k_KQ_0 += WARP_SIZE) { | ~~~~~~ ^ ~~~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:490:39: note: in instantiation of function template specialization 'vec_dot_fattn_vec_KQ_q4_0' requested here 490 | return type_K == GGML_TYPE_Q4_0 ? vec_dot_fattn_vec_KQ_q4_0 : | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:318:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 318 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:130:17: note: in instantiation of function template specialization 'quantize_q8_1_to_shared>' requested here 130 | quantize_q8_1_to_shared(Q_f + 4*i0, scale, tmp_q_i32, tmp_q_ds); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:284:35: note: in instantiation of function template specialization 'flash_attn_vec_ext_f32<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 284 | fattn_kernel_t fattn_kernel = flash_attn_vec_ext_f32; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:325:23: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 325 | for (int l = 1; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:341:27: warning: comparison of integers of different signs: 'int' and 'unsigned long' [-Wsign-compare] 341 | for (int l = 0; l < sizeof(int); ++l) { | ~ ^ ~~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 4>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 4>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:308:13: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 1, 4, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 308 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:2: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:554:24: warning: cast from 'const float *' to 'unsigned int *' drops const qualifier [-Wcast-qual] 554 | *((uint32_t *) &KQ_max_scale) &= ftz_mask; | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-common.cuh:704:5: note: in instantiation of function template specialization 'flash_attn_combine_results<256, 1>' requested here 704 | flash_attn_combine_results | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:287:5: note: in instantiation of function template specialization 'launch_fattn<256, 1>' requested here 287 | launch_fattn(ctx, dst, fattn_kernel, nwarps, cols_per_block, need_f16_K, need_f16_V); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../fattn-vec-f32.cuh:359:9: note: in instantiation of function template specialization 'ggml_cuda_flash_attn_ext_vec_f32_case_impl<256, 8, 1, GGML_TYPE_F16, GGML_TYPE_F16, false>' requested here 359 | ggml_cuda_flash_attn_ext_vec_f32_case_impl(ctx, dst); | ^ 34 warnings generated when compiling for host. 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q2_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. 78 warnings generated when compiling for gfx942. In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../common.cuh:20: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:171:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 171 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:192:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 192 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:213:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 213 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:254:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 254 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:281:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 281 | struct { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-hip/../ggml-common.h:298:9: warning: anonymous types declared in an anonymous union are an extension [-Wnested-anon-types] 298 | struct { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:5: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:172:105: warning: function 'mma_K4' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 172 | __device__ __forceinline__ void mma_K4(const mma_int_A_I16K4 & mma_A, const mma_int_B_J8K4 & mma_B) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mma.cuh:194:105: warning: function 'mma_K8' could be declared with attribute 'noreturn' [-Wmissing-noreturn] 194 | __device__ __forceinline__ void mma_K8(const mma_int_A_I16K8 & mma_A, const mma_int_B_J8K8 & mma_B) { | ^ In file included from /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/mmq-instance-q4_k.cu:3: /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:866:99: warning: unused parameter 'k00' [-Wunused-parameter] 866 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1053:99: warning: unused parameter 'k00' [-Wunused-parameter] 1053 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:1704:99: warning: unused parameter 'k00' [-Wunused-parameter] 1704 | const int * __restrict__ x, const int * __restrict__ y, float * __restrict__ sum, const int & k00) { | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:17: warning: unused parameter 'ne00' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2497:75: warning: unused parameter 'ne10' [-Wunused-parameter] 2497 | const int & ne00, const int & ne01, const int & stride01, const int & ne10, const int & ne11, const int & stride11, const int & ne0, | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2821:15: warning: unused variable 'nsm' [-Wunused-variable] 2821 | const int nsm = ggml_cuda_info().devices[id].nsm; | ^~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2852:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2852 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2855:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2855 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2858:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2858 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2861:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2861 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2864:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2864 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2867:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2867 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2870:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2870 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2873:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2873 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2876:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2876 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2879:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2879 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2882:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2882 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2885:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2885 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2888:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2888 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2891:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2891 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2894:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2894 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2805:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2805 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:36: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2813:9: note: in instantiation of function template specialization 'mul_mat_q_stream_k_fixup' requested here 2813 | mul_mat_q_stream_k_fixup<<>> | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2897:13: note: in instantiation of function template specialization 'launch_mul_mat_q' requested here 2897 | launch_mul_mat_q(ctx, args, stream); | ^ /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-cuda/template-instances/../mmq.cuh:2691:16: warning: comparison of integers of different signs: 'const int' and 'unsigned int' [-Wsign-compare] 2691 | if (it != blockIdx.x || jt != blockIdx.y) { | ~~ ^ ~~~~~~~~~~ 78 warnings generated when compiling for host. [ 75%] Linking CXX shared library ../../../bin/libggml-hip.so cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/cmake -E cmake_link_script CMakeFiles/ggml-hip.dir/link.txt --verbose=1 /usr/bin/hipcc -fPIC -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -Xlinker --dependency-file=CMakeFiles/ggml-hip.dir/link.d -Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes -shared -Wl,-soname,libggml-hip.so.b4580 -o ../../../bin/libggml-hip.so.b4580 "CMakeFiles/ggml-hip.dir/__/ggml-cuda/acc.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/arange.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/argmax.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/argsort.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/binbcast.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/clamp.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/concat.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/conv-transpose-1d.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/convert.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/count-equal.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/cpy.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/cross-entropy-loss.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/diagmask.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn-tile-f32.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/fattn.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/getrows.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/ggml-cuda.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/gla.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/im2col.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmq.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmv.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/mmvq.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/norm.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/opt-step-adamw.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/out-prod.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/pad.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/pool2d.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/quantize.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/rope.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/scale.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/softmax.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/sum.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/sumrows.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/tsembd.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/unary.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/upscale.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/wkv6.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqfloat-cpb32.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb32.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-wmma-f16-instance-kqhalf-cpb8.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq1_s.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_s.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xs.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq2_xxs.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_s.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq3_xxs.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_nl.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-iq4_xs.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q2_k.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q3_k.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-iclang++: warning: argument unused during compilation: '-Xarch_host -fstack-protector-strong' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-Xarch_host -fcf-protection' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-specs=/usr/lib/rpm/redhat/redhat-package-notes' [-Wunused-command-line-argument] nstance-q4_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_1.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q4_k.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_1.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q5_k.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q6_k.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/mmq-instance-q8_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q4_0-q4_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q4_0-q4_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-q8_0-q8_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-q8_0-q8_0.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs128-f16-f16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs256-f16-f16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f16-instance-hs64-f16-f16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs128-f16-f16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs256-f16-f16.cu.o" "CMakeFiles/ggml-hip.dir/__/ggml-cuda/template-instances/fattn-vec-f32-instance-hs64-f16-f16.cu.o" ../../../bin/libggml-base.so.b4580 /usr/lib64/libhipblas.so.2.4 --hip-link --offload-arch=gfx900 --offload-arch=gfx906:xnack- --offload-arch=gfx908:xnack- --offload-arch=gfx90a:xnack+ --offload-arch=gfx90a:xnack- --offload-arch=gfx942 --offload-arch=gfx1010 --offload-arch=gfx1012 --offload-arch=gfx1030 --offload-arch=gfx1031 --offload-arch=gfx1035 --offload-arch=gfx1100 --offload-arch=gfx1101 --offload-arch=gfx1102 --offload-arch=gfx1103 --offload-arch=gfx1150 --offload-arch=gfx1151 --offload-arch=gfx1152 --offload-arch=gfx1200 --offload-arch=gfx1201 /usr/lib64/librocblas.so.4.4 /usr/lib64/libamdhip64.so.6.4.43482 cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/ggml-hip && /usr/bin/cmake -E cmake_symlink_library ../../../bin/libggml-hip.so.b4580 ../../../bin/libggml-hip.so.b4580 ../../../bin/libggml-hip.so gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 75%] Built target ggml-hip /usr/bin/gmake -f ggml/src/CMakeFiles/ggml.dir/build.make ggml/src/CMakeFiles/ggml.dir/depend gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src/CMakeFiles/ggml.dir/DependInfo.cmake "--color=" gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f ggml/src/CMakeFiles/ggml.dir/build.make ggml/src/CMakeFiles/ggml.dir/build gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 75%] Building CXX object ggml/src/CMakeFiles/ggml.dir/ggml-backend-reg.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_BUILD -DGGML_SCHED_MAX_COPIES=4 -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -D_GNU_SOURCE -D_XOPEN_SOURCE=600 -Dggml_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -std=gnu++17 -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT ggml/src/CMakeFiles/ggml.dir/ggml-backend-reg.cpp.o -MF CMakeFiles/ggml.dir/ggml-backend-reg.cpp.o.d -o CMakeFiles/ggml.dir/ggml-backend-reg.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/ggml-backend-reg.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 76%] Linking CXX shared library ../../bin/libggml.so cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/cmake -E cmake_link_script CMakeFiles/ggml.dir/link.txt --verbose=1 sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory clang++: warning: argument unused during compilation: '-Xarch_host -fstack-protector-strong' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-Xarch_host -fcf-protection' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-specs=/usr/lib/rpm/redhat/redhat-package-notes' [-Wunused-command-line-argument] /usr/bin/hipcc -fPIC -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -Xlinker --dependency-file=CMakeFiles/ggml.dir/link.d -Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes -shared -Wl,-soname,libggml.so.b4580 -o ../../bin/libggml.so.b4580 "CMakeFiles/ggml.dir/ggml-backend-reg.cpp.o" -ldl ../../bin/libggml-cpu.so.b4580 ../../bin/libggml-hip.so.b4580 ../../bin/libggml-base.so.b4580 cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/ggml/src && /usr/bin/cmake -E cmake_symlink_library ../../bin/libggml.so.b4580 ../../bin/libggml.so.b4580 ../../bin/libggml.so gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 76%] Built target ggml /usr/bin/gmake -f src/CMakeFiles/llama.dir/build.make src/CMakeFiles/llama.dir/depend gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src/CMakeFiles/llama.dir/DependInfo.cmake "--color=" gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f src/CMakeFiles/llama.dir/build.make src/CMakeFiles/llama.dir/build gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 77%] Building CXX object src/CMakeFiles/llama.dir/llama.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama.cpp.o -MF CMakeFiles/llama.dir/llama.cpp.o.d -o CMakeFiles/llama.dir/llama.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama.cpp [ 78%] Building CXX object src/CMakeFiles/llama.dir/llama-context.cpp.o [ 79%] Building CXX object src/CMakeFiles/llama.dir/llama-batch.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-context.cpp.o -MF CMakeFiles/llama.dir/llama-context.cpp.o.d -o CMakeFiles/llama.dir/llama-context.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-context.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-batch.cpp.o -MF CMakeFiles/llama.dir/llama-batch.cpp.o.d -o CMakeFiles/llama.dir/llama-batch.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-batch.cpp [ 79%] Building CXX object src/CMakeFiles/llama.dir/llama-arch.cpp.o [ 80%] Building CXX object src/CMakeFiles/llama.dir/llama-adapter.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-arch.cpp.o -MF CMakeFiles/llama.dir/llama-arch.cpp.o.d -o CMakeFiles/llama.dir/llama-arch.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-arch.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-adapter.cpp.o -MF CMakeFiles/llama.dir/llama-adapter.cpp.o.d -o CMakeFiles/llama.dir/llama-adapter.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-adapter.cpp [ 81%] Building CXX object src/CMakeFiles/llama.dir/llama-chat.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-chat.cpp.o -MF CMakeFiles/llama.dir/llama-chat.cpp.o.d -o CMakeFiles/llama.dir/llama-chat.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-chat.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 82%] Building CXX object src/CMakeFiles/llama.dir/llama-grammar.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-grammar.cpp.o -MF CMakeFiles/llama.dir/llama-grammar.cpp.o.d -o CMakeFiles/llama.dir/llama-grammar.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-grammar.cpp [ 82%] Building CXX object src/CMakeFiles/llama.dir/llama-hparams.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-hparams.cpp.o -MF CMakeFiles/llama.dir/llama-hparams.cpp.o.d -o CMakeFiles/llama.dir/llama-hparams.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-hparams.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 83%] Building CXX object src/CMakeFiles/llama.dir/llama-impl.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-impl.cpp.o -MF CMakeFiles/llama.dir/llama-impl.cpp.o.d -o CMakeFiles/llama.dir/llama-impl.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-impl.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 84%] Building CXX object src/CMakeFiles/llama.dir/llama-kv-cache.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-kv-cache.cpp.o -MF CMakeFiles/llama.dir/llama-kv-cache.cpp.o.d -o CMakeFiles/llama.dir/llama-kv-cache.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-kv-cache.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 85%] Building CXX object src/CMakeFiles/llama.dir/llama-mmap.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-mmap.cpp.o -MF CMakeFiles/llama.dir/llama-mmap.cpp.o.d -o CMakeFiles/llama.dir/llama-mmap.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-mmap.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 86%] Building CXX object src/CMakeFiles/llama.dir/llama-model-loader.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-model-loader.cpp.o -MF CMakeFiles/llama.dir/llama-model-loader.cpp.o.d -o CMakeFiles/llama.dir/llama-model-loader.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-model-loader.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 87%] Building CXX object src/CMakeFiles/llama.dir/llama-model.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-model.cpp.o -MF CMakeFiles/llama.dir/llama-model.cpp.o.d -o CMakeFiles/llama.dir/llama-model.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-model.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 87%] Building CXX object src/CMakeFiles/llama.dir/llama-quant.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-quant.cpp.o -MF CMakeFiles/llama.dir/llama-quant.cpp.o.d -o CMakeFiles/llama.dir/llama-quant.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-quant.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 88%] Building CXX object src/CMakeFiles/llama.dir/llama-sampling.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-sampling.cpp.o -MF CMakeFiles/llama.dir/llama-sampling.cpp.o.d -o CMakeFiles/llama.dir/llama-sampling.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-sampling.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 89%] Building CXX object src/CMakeFiles/llama.dir/llama-vocab.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/llama-vocab.cpp.o -MF CMakeFiles/llama.dir/llama-vocab.cpp.o.d -o CMakeFiles/llama.dir/llama-vocab.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/llama-vocab.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 90%] Building CXX object src/CMakeFiles/llama.dir/unicode.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/unicode.cpp.o -MF CMakeFiles/llama.dir/unicode.cpp.o.d -o CMakeFiles/llama.dir/unicode.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/unicode.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 91%] Building CXX object src/CMakeFiles/llama.dir/unicode-data.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_BUILD -DLLAMA_SHARED -Dllama_EXPORTS -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT src/CMakeFiles/llama.dir/unicode-data.cpp.o -MF CMakeFiles/llama.dir/unicode-data.cpp.o.d -o CMakeFiles/llama.dir/unicode-data.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/unicode-data.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 92%] Linking CXX shared library ../bin/libllama.so cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/cmake -E cmake_link_script CMakeFiles/llama.dir/link.txt --verbose=1 sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory clang++: warning: argument unused during compilation: '-Xarch_host -fstack-protector-strong' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-Xarch_host -fcf-protection' [-Wunused-command-line-argument] clang++: warning: argument unused during compilation: '-specs=/usr/lib/rpm/redhat/redhat-package-notes' [-Wunused-command-line-argument] /usr/bin/hipcc -fPIC -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -Xlinker --dependency-file=CMakeFiles/llama.dir/link.d -Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes -shared -Wl,-soname,libllama.so.b4580 -o ../bin/libllama.so.b4580 CMakeFiles/llama.dir/llama.cpp.o "CMakeFiles/llama.dir/llama-adapter.cpp.o" "CMakeFiles/llama.dir/llama-arch.cpp.o" "CMakeFiles/llama.dir/llama-batch.cpp.o" "CMakeFiles/llama.dir/llama-chat.cpp.o" "CMakeFiles/llama.dir/llama-context.cpp.o" "CMakeFiles/llama.dir/llama-grammar.cpp.o" "CMakeFiles/llama.dir/llama-hparams.cpp.o" "CMakeFiles/llama.dir/llama-impl.cpp.o" "CMakeFiles/llama.dir/llama-kv-cache.cpp.o" "CMakeFiles/llama.dir/llama-mmap.cpp.o" "CMakeFiles/llama.dir/llama-model-loader.cpp.o" "CMakeFiles/llama.dir/llama-model.cpp.o" "CMakeFiles/llama.dir/llama-quant.cpp.o" "CMakeFiles/llama.dir/llama-sampling.cpp.o" "CMakeFiles/llama.dir/llama-vocab.cpp.o" CMakeFiles/llama.dir/unicode.cpp.o "CMakeFiles/llama.dir/unicode-data.cpp.o" ../bin/libggml.so.b4580 ../bin/libggml-cpu.so.b4580 ../bin/libggml-hip.so.b4580 ../bin/libggml-base.so.b4580 cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/src && /usr/bin/cmake -E cmake_symlink_library ../bin/libllama.so.b4580 ../bin/libllama.so.b4580 ../bin/libllama.so gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 92%] Built target llama /usr/bin/gmake -f common/CMakeFiles/common.dir/build.make common/CMakeFiles/common.dir/depend gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build && /usr/bin/cmake -E cmake_depends "Unix Makefiles" /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common/CMakeFiles/common.dir/DependInfo.cmake "--color=" gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/gmake -f common/CMakeFiles/common.dir/build.make common/CMakeFiles/common.dir/build gmake[2]: Entering directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [ 94%] Building CXX object common/CMakeFiles/common.dir/arg.cpp.o [ 95%] Building CXX object common/CMakeFiles/common.dir/json-schema-to-grammar.cpp.o [ 95%] Building CXX object common/CMakeFiles/common.dir/console.cpp.o [ 95%] Building CXX object common/CMakeFiles/common.dir/common.cpp.o [ 96%] Building CXX object common/CMakeFiles/common.dir/ngram-cache.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/arg.cpp.o -MF CMakeFiles/common.dir/arg.cpp.o.d -o CMakeFiles/common.dir/arg.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/arg.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/common.cpp.o -MF CMakeFiles/common.dir/common.cpp.o.d -o CMakeFiles/common.dir/common.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/common.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/console.cpp.o -MF CMakeFiles/common.dir/console.cpp.o.d -o CMakeFiles/common.dir/console.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/console.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/json-schema-to-grammar.cpp.o -MF CMakeFiles/common.dir/json-schema-to-grammar.cpp.o.d -o CMakeFiles/common.dir/json-schema-to-grammar.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/json-schema-to-grammar.cpp cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/ngram-cache.cpp.o -MF CMakeFiles/common.dir/ngram-cache.cpp.o.d -o CMakeFiles/common.dir/ngram-cache.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/ngram-cache.cpp [ 97%] Building CXX object common/CMakeFiles/common.dir/log.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/log.cpp.o -MF CMakeFiles/common.dir/log.cpp.o.d -o CMakeFiles/common.dir/log.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/log.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 98%] Building CXX object common/CMakeFiles/common.dir/sampling.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/sampling.cpp.o -MF CMakeFiles/common.dir/sampling.cpp.o.d -o CMakeFiles/common.dir/sampling.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/sampling.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [ 99%] Building CXX object common/CMakeFiles/common.dir/speculative.cpp.o cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/hipcc -DGGML_BACKEND_SHARED -DGGML_SHARED -DGGML_USE_CPU -DGGML_USE_CUDA -DGGML_USE_HIP -DLLAMA_SHARED -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/. -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../include -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/src/../common -I/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/ggml/src/../include -O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -DNDEBUG -fPIC -Wmissing-declarations -Wmissing-noreturn -Wall -Wextra -Wpedantic -Wcast-qual -Wno-unused-function -Wunreachable-code-break -Wunreachable-code-return -Wmissing-prototypes -Wextra-semi -MD -MT common/CMakeFiles/common.dir/speculative.cpp.o -MF CMakeFiles/common.dir/speculative.cpp.o.d -o CMakeFiles/common.dir/speculative.cpp.o -c /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/common/speculative.cpp sh: line 1: /usr/bin/rocm_agent_enumerator: No such file or directory [100%] Linking CXX static library libcommon.a cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/cmake -P CMakeFiles/common.dir/cmake_clean_target.cmake cd /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/common && /usr/bin/cmake -E cmake_link_script CMakeFiles/common.dir/link.txt --verbose=1 bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record bfd plugin: LLVM gold plugin has failed to create LTO module: Invalid record /usr/bin/ar qc libcommon.a CMakeFiles/common.dir/arg.cpp.o CMakeFiles/common.dir/common.cpp.o CMakeFiles/common.dir/console.cpp.o "CMakeFiles/common.dir/json-schema-to-grammar.cpp.o" CMakeFiles/common.dir/log.cpp.o "CMakeFiles/common.dir/ngram-cache.cpp.o" CMakeFiles/common.dir/sampling.cpp.o CMakeFiles/common.dir/speculative.cpp.o "CMakeFiles/build_info.dir/build-info.cpp.o" /usr/bin/ranlib libcommon.a gmake[2]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' [100%] Built target common gmake[1]: Leaving directory '/builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build' /usr/bin/cmake -E cmake_progress_start /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/redhat-linux-build/CMakeFiles 0 + RPM_EC=0 ++ jobs -p + exit 0 Executing(%install): /bin/sh -e /var/tmp/rpm-tmp.B4jFdU + umask 022 + cd /builddir/build/BUILD/llama-cpp-b4580-build + '[' /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT '!=' / ']' + rm -rf /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT ++ dirname /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT + mkdir -p /builddir/build/BUILD/llama-cpp-b4580-build + mkdir /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT + CFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer ' + export CFLAGS + CXXFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Werror=format-security -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -Xarch_host -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -Xarch_host -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer' + export CXXFLAGS + FFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -I/usr/lib64/gfortran/modules ' + export FFLAGS + FCFLAGS='-O2 -flto=thin -fexceptions -g -grecord-gcc-switches -pipe -Wall -Wp,-U_FORTIFY_SOURCE,-D_FORTIFY_SOURCE=3 -Wp,-D_GLIBCXX_ASSERTIONS --config /usr/lib/rpm/redhat/redhat-hardened-clang.cfg -fstack-protector-strong -m64 -march=x86-64 -mtune=generic -fasynchronous-unwind-tables -fstack-clash-protection -fcf-protection -fno-omit-frame-pointer -mno-omit-leaf-frame-pointer -I/usr/lib64/gfortran/modules ' + export FCFLAGS + VALAFLAGS=-g + export VALAFLAGS + RUSTFLAGS='-Copt-level=3 -Cdebuginfo=2 -Ccodegen-units=1 -Cstrip=none -Cforce-frame-pointers=yes -Clink-arg=-specs=/usr/lib/rpm/redhat/redhat-package-notes --cap-lints=warn' + export RUSTFLAGS + LDFLAGS='-Wl,-z,relro -Wl,--as-needed -Wl,-z,pack-relative-relocs -Wl,-z,now -Wl,-z,now -Wl,--build-id=sha1 -specs=/usr/lib/rpm/redhat/redhat-package-notes ' + export LDFLAGS + LT_SYS_LIBRARY_PATH=/usr/lib64: + export LT_SYS_LIBRARY_PATH + CC=hipcc + export CC + CXX=hipcc + export CXX + cd llama.cpp-b4580 + DESTDIR=/builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT + /usr/bin/cmake --install redhat-linux-build -- Install configuration: "Release" -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml-cpu.so.b4580 -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml-cpu.so -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml-hip.so.b4580 -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml-hip.so -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml.so.b4580 -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml.so -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-cpu.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-alloc.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-backend.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-blas.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-cann.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-cuda.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-kompute.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-opt.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-metal.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-rpc.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-sycl.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/ggml-vulkan.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/gguf.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml-base.so.b4580 -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml-base.so -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/cmake/ggml/ggml-config.cmake -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/cmake/ggml/ggml-version.cmake -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libllama.so.b4580 -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libllama.so -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/llama.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/include/llama-cpp.h -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/cmake/llama/llama-config.cmake -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/cmake/llama/llama-version.cmake -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/bin/convert_hf_to_gguf.py -- Installing: /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib/pkgconfig/llama.pc + rm -rf '/builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/lib64/libggml_shared.*' + rm /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/bin/convert_hf_to_gguf.py + /usr/bin/find-debuginfo -j6 --strict-build-id -m -i --build-id-seed b4580-2.fc43 --unique-debug-suffix -b4580-2.fc43.x86_64 --unique-debug-src-base llama-cpp-b4580-2.fc43.x86_64 --run-dwz --dwz-low-mem-die-limit 10000000 --dwz-max-die-limit 110000000 -S debugsourcefiles.list /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580 find-debuginfo: starting Extracting debug info from 5 files DWARF-compressing 5 files dwz: ./usr/lib64/libggml-base.so.b4580-b4580-2.fc43.x86_64.debug: Unknown debugging section .debug_str_offsets dwz: ./usr/lib64/libggml-cpu.so.b4580-b4580-2.fc43.x86_64.debug: Unknown debugging section .debug_str_offsets dwz: ./usr/lib64/libggml-hip.so.b4580-b4580-2.fc43.x86_64.debug: Unknown debugging section .debug_str_offsets dwz: ./usr/lib64/libggml.so.b4580-b4580-2.fc43.x86_64.debug: Unknown debugging section .debug_str_offsets dwz: ./usr/lib64/libllama.so.b4580-b4580-2.fc43.x86_64.debug: Unknown debugging section .debug_str_offsets dwz: Too few files for multifile optimization sepdebugcrcfix: Updated 0 CRC32s, 5 CRC32s did match. Creating .debug symlinks for symlinks to ELF files Copying sources found by 'debugedit -l' to /usr/src/debug/llama-cpp-b4580-2.fc43.x86_64 find-debuginfo: done + /usr/lib/rpm/check-buildroot + /usr/lib/rpm/redhat/brp-ldconfig + /usr/lib/rpm/brp-compress + /usr/lib/rpm/redhat/brp-strip-lto /usr/bin/strip + /usr/lib/rpm/brp-strip-static-archive /usr/bin/strip + /usr/lib/rpm/check-rpaths + /usr/lib/rpm/redhat/brp-mangle-shebangs + /usr/lib/rpm/brp-remove-la-files + /usr/lib/rpm/redhat/brp-python-rpm-in-distinfo + env /usr/lib/rpm/redhat/brp-python-bytecompile '' 1 0 -j6 + /usr/lib/rpm/redhat/brp-python-hardlink + /usr/bin/add-determinism --brp -j6 /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT Scanned 32 directories and 208 files, processed 0 inodes, 0 modified (0 replaced + 0 rewritten), 0 unsupported format, 0 errors Reading /builddir/build/BUILD/llama-cpp-b4580-build/SPECPARTS/rpm-debuginfo.specpart Processing files: llama-cpp-b4580-2.fc43.x86_64 Executing(%license): /bin/sh -e /var/tmp/rpm-tmp.FoFP36 + umask 022 + cd /builddir/build/BUILD/llama-cpp-b4580-build + cd llama.cpp-b4580 + LICENSEDIR=/builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/share/licenses/llama-cpp + export LC_ALL=C.UTF-8 + LC_ALL=C.UTF-8 + export LICENSEDIR + /usr/bin/mkdir -p /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/share/licenses/llama-cpp + cp -pr /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/LICENSE /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/share/licenses/llama-cpp + RPM_EC=0 ++ jobs -p + exit 0 Provides: libggml-base.so.b4580()(64bit) libggml-cpu.so.b4580()(64bit) libggml-hip.so.b4580()(64bit) libggml.so.b4580()(64bit) libllama.so.b4580()(64bit) llama-cpp = b4580-2.fc43 llama-cpp(x86-64) = b4580-2.fc43 Requires(rpmlib): rpmlib(CompressedFileNames) <= 3.0.4-1 rpmlib(FileDigests) <= 4.6.0-1 rpmlib(PayloadFilesHavePrefix) <= 4.0-1 Requires: ld-linux-x86-64.so.2()(64bit) ld-linux-x86-64.so.2(GLIBC_2.3)(64bit) libamdhip64.so.6()(64bit) libamdhip64.so.6(hip_4.2)(64bit) libamdhip64.so.6(hip_6.0)(64bit) libc.so.6()(64bit) libc.so.6(GLIBC_2.14)(64bit) libc.so.6(GLIBC_2.17)(64bit) libc.so.6(GLIBC_2.2.5)(64bit) libc.so.6(GLIBC_2.29)(64bit) libc.so.6(GLIBC_2.3.2)(64bit) libc.so.6(GLIBC_2.3.4)(64bit) libc.so.6(GLIBC_2.32)(64bit) libc.so.6(GLIBC_2.33)(64bit) libc.so.6(GLIBC_2.34)(64bit) libc.so.6(GLIBC_2.38)(64bit) libc.so.6(GLIBC_2.4)(64bit) libc.so.6(GLIBC_2.7)(64bit) libc.so.6(GLIBC_ABI_DT_RELR)(64bit) libgcc_s.so.1()(64bit) libgcc_s.so.1(GCC_3.0)(64bit) libggml-base.so.b4580()(64bit) libggml-cpu.so.b4580()(64bit) libggml-hip.so.b4580()(64bit) libggml.so.b4580()(64bit) libhipblas.so.2()(64bit) libm.so.6()(64bit) libm.so.6(GLIBC_2.2.5)(64bit) libm.so.6(GLIBC_2.27)(64bit) libm.so.6(GLIBC_2.29)(64bit) librocblas.so.4()(64bit) libstdc++.so.6()(64bit) libstdc++.so.6(CXXABI_1.3)(64bit) libstdc++.so.6(CXXABI_1.3.11)(64bit) libstdc++.so.6(CXXABI_1.3.13)(64bit) libstdc++.so.6(CXXABI_1.3.2)(64bit) libstdc++.so.6(CXXABI_1.3.3)(64bit) libstdc++.so.6(CXXABI_1.3.5)(64bit) libstdc++.so.6(CXXABI_1.3.9)(64bit) libstdc++.so.6(GLIBCXX_3.4)(64bit) libstdc++.so.6(GLIBCXX_3.4.11)(64bit) libstdc++.so.6(GLIBCXX_3.4.14)(64bit) libstdc++.so.6(GLIBCXX_3.4.15)(64bit) libstdc++.so.6(GLIBCXX_3.4.17)(64bit) libstdc++.so.6(GLIBCXX_3.4.18)(64bit) libstdc++.so.6(GLIBCXX_3.4.19)(64bit) libstdc++.so.6(GLIBCXX_3.4.20)(64bit) libstdc++.so.6(GLIBCXX_3.4.21)(64bit) libstdc++.so.6(GLIBCXX_3.4.22)(64bit) libstdc++.so.6(GLIBCXX_3.4.25)(64bit) libstdc++.so.6(GLIBCXX_3.4.26)(64bit) libstdc++.so.6(GLIBCXX_3.4.29)(64bit) libstdc++.so.6(GLIBCXX_3.4.30)(64bit) libstdc++.so.6(GLIBCXX_3.4.9)(64bit) Recommends: numactl Processing files: llama-cpp-devel-b4580-2.fc43.x86_64 Executing(%doc): /bin/sh -e /var/tmp/rpm-tmp.gccnHa + umask 022 + cd /builddir/build/BUILD/llama-cpp-b4580-build + cd llama.cpp-b4580 + DOCDIR=/builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/share/doc/llama-cpp-devel + export LC_ALL=C.UTF-8 + LC_ALL=C.UTF-8 + export DOCDIR + /usr/bin/mkdir -p /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/share/doc/llama-cpp-devel + cp -pr /builddir/build/BUILD/llama-cpp-b4580-build/llama.cpp-b4580/README.md /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT/usr/share/doc/llama-cpp-devel + RPM_EC=0 ++ jobs -p + exit 0 Provides: cmake(ggml) cmake(llama) llama-cpp-devel = b4580-2.fc43 llama-cpp-devel(x86-64) = b4580-2.fc43 Requires(rpmlib): rpmlib(CompressedFileNames) <= 3.0.4-1 rpmlib(FileDigests) <= 4.6.0-1 rpmlib(PayloadFilesHavePrefix) <= 4.0-1 Requires: cmake-filesystem(x86-64) libggml-base.so.b4580()(64bit) libggml-cpu.so.b4580()(64bit) libggml-hip.so.b4580()(64bit) libggml.so.b4580()(64bit) libllama.so.b4580()(64bit) Processing files: llama-cpp-debugsource-b4580-2.fc43.x86_64 Provides: llama-cpp-debugsource = b4580-2.fc43 llama-cpp-debugsource(x86-64) = b4580-2.fc43 Requires(rpmlib): rpmlib(CompressedFileNames) <= 3.0.4-1 rpmlib(FileDigests) <= 4.6.0-1 rpmlib(PayloadFilesHavePrefix) <= 4.0-1 Processing files: llama-cpp-debuginfo-b4580-2.fc43.x86_64 Provides: debuginfo(build-id) = 1717a0b1f3f3b47f5723aae9ecbadc7e896a2a38 debuginfo(build-id) = d3cc48a3378229a4f5fab9f17310acb1fe0a30ef debuginfo(build-id) = ddfdb80174293e4b6eab0d9e1aab752e03102d4e debuginfo(build-id) = ed24474c415112561475acdd4541217524eb6e1b debuginfo(build-id) = f75d19df4cec2ed3e24552a87b814ea3cd18c420 libggml-base.so.b4580-b4580-2.fc43.x86_64.debug()(64bit) libggml-cpu.so.b4580-b4580-2.fc43.x86_64.debug()(64bit) libggml-hip.so.b4580-b4580-2.fc43.x86_64.debug()(64bit) libggml.so.b4580-b4580-2.fc43.x86_64.debug()(64bit) libllama.so.b4580-b4580-2.fc43.x86_64.debug()(64bit) llama-cpp-debuginfo = b4580-2.fc43 llama-cpp-debuginfo(x86-64) = b4580-2.fc43 Requires(rpmlib): rpmlib(CompressedFileNames) <= 3.0.4-1 rpmlib(FileDigests) <= 4.6.0-1 rpmlib(PayloadFilesHavePrefix) <= 4.0-1 Recommends: llama-cpp-debugsource(x86-64) = b4580-2.fc43 Checking for unpackaged file(s): /usr/lib/rpm/check-files /builddir/build/BUILD/llama-cpp-b4580-build/BUILDROOT Wrote: /builddir/build/RPMS/llama-cpp-devel-b4580-2.fc43.x86_64.rpm Wrote: /builddir/build/RPMS/llama-cpp-debugsource-b4580-2.fc43.x86_64.rpm Wrote: /builddir/build/RPMS/llama-cpp-debuginfo-b4580-2.fc43.x86_64.rpm Wrote: /builddir/build/RPMS/llama-cpp-b4580-2.fc43.x86_64.rpm Child return code was: 0